state: generation 13283 (cycle)

This commit is contained in:
2026-09-23 02:15:32 +00:00
parent 6149e3410c
commit 6c8149ad25
10 changed files with 1879 additions and 1911 deletions

View File

@@ -3044,7 +3044,7 @@
"taskType": "text-generation" "taskType": "text-generation"
} }
}, },
"generatedAt": "2026-09-23T02:10:36.863304+00:00", "generatedAt": "2026-09-23T02:15:14.561131+00:00",
"summary": { "summary": {
"activeBlockCount": 153, "activeBlockCount": 153,
"byGpuFramework": { "byGpuFramework": {

View File

@@ -396,7 +396,7 @@
} }
}, },
"frameworkUpdatedAt": null, "frameworkUpdatedAt": null,
"generatedAt": "2026-09-23T02:14:11.149033+00:00", "generatedAt": "2026-09-23T02:15:31.588052+00:00",
"gpuStats": { "gpuStats": {
"Ascend_910-b4": { "Ascend_910-b4": {
"available": true, "available": true,

View File

@@ -1,5 +1,5 @@
{ {
"catalogUpdatedAt": "2026-09-23T02:14:11.149033+00:00", "catalogUpdatedAt": "2026-09-23T02:15:31.588052+00:00",
"configuredTaskTypes": [ "configuredTaskTypes": [
"text-generation" "text-generation"
], ],
@@ -56,7 +56,7 @@
"time-series-forecasting" "time-series-forecasting"
], ],
"errors": [], "errors": [],
"generatedAt": "2026-09-23T02:14:11.571364+00:00", "generatedAt": "2026-09-23T02:15:31.588052+00:00",
"gpuCatalog": { "gpuCatalog": {
"Ascend_910-b3": { "Ascend_910-b3": {
"canVerify": true, "canVerify": true,
@@ -6669,6 +6669,6 @@
"updateTime": "2025-12-22 08:59:53" "updateTime": "2025-12-22 08:59:53"
} }
], ],
"taskTreeUpdatedAt": "2026-09-23T02:14:11.149033+00:00", "taskTreeUpdatedAt": "2026-09-23T02:15:31.588052+00:00",
"version": 1 "version": 1
} }

View File

@@ -1,6 +1,6 @@
{ {
"generatedAt": "2026-09-23T02:06:47.498515+00:00", "generatedAt": "2026-09-23T02:15:14.481703+00:00",
"lastSyncTime": "2026-09-23T02:06:44.350388+00:00", "lastSyncTime": "2026-09-23T02:15:13.670755+00:00",
"recentLimit": 300, "recentLimit": 300,
"report": { "report": {
"architectureCompatibilityBlocks": { "architectureCompatibilityBlocks": {
@@ -5158,28 +5158,28 @@
"unresolvedFailureCount": 296 "unresolvedFailureCount": 296
}, },
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": { "Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": {
"attributableFailureCount": 102, "attributableFailureCount": 103,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 102, "decisionTotal": 103,
"failureBreakdown": { "failureBreakdown": {
"backend_operator": 8, "backend_operator": 9,
"framework_architecture_unsupported": 12, "framework_architecture_unsupported": 12,
"platform_infrastructure": 18, "platform_infrastructure": 19,
"tokenizer_compatibility": 82, "tokenizer_compatibility": 82,
"参数/模板问题": 6 "参数/模板问题": 6
}, },
"failureCount": 126, "failureCount": 128,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm_fix_tokenizer", "framework": "vllm_fix_tokenizer",
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 18, "platformFailureCount": 19,
"successCount": 0, "successCount": 0,
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1", "targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation", "taskType": "text-generation",
"total": 126, "total": 128,
"unresolvedFailureCount": 6 "unresolvedFailureCount": 6
}, },
"Sunrise_pt-200-x1|vllm|text-generation": { "Sunrise_pt-200-x1|vllm|text-generation": {
@@ -5835,28 +5835,28 @@
"unresolvedFailureCount": 25 "unresolvedFailureCount": 25
}, },
"vllm_fix_tokenizer": { "vllm_fix_tokenizer": {
"attributableFailureCount": 134, "attributableFailureCount": 135,
"decisionFailureRate": 0.964, "decisionFailureRate": 0.9643,
"decisionSuccessRate": 0.036, "decisionSuccessRate": 0.0357,
"decisionTotal": 139, "decisionTotal": 140,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 298, "ambiguous_runtime": 298,
"backend_operator": 8, "backend_operator": 9,
"framework_architecture_unsupported": 29, "framework_architecture_unsupported": 29,
"memory_capacity": 10, "memory_capacity": 10,
"model_load": 5, "model_load": 5,
"platform_infrastructure": 18, "platform_infrastructure": 19,
"tokenizer_compatibility": 82, "tokenizer_compatibility": 82,
"参数/模板问题": 13 "参数/模板问题": 13
}, },
"failureCount": 463, "failureCount": 465,
"failureRate": 0.9893, "failureRate": 0.9894,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 18, "platformFailureCount": 19,
"successCount": 5, "successCount": 5,
"successRate": 0.0107, "successRate": 0.0106,
"total": 468, "total": 470,
"unresolvedFailureCount": 311 "unresolvedFailureCount": 311
}, },
"vllm_tokenizer_patch": { "vllm_tokenizer_patch": {
@@ -5888,7 +5888,7 @@
"unresolvedFailureCount": 175 "unresolvedFailureCount": 175
} }
}, },
"generatedAt": "2026-09-23T02:06:47.486688+00:00", "generatedAt": "2026-09-23T02:15:14.468487+00:00",
"gpuSummaries": { "gpuSummaries": {
"Ascend_910-b3": { "Ascend_910-b3": {
"attributableFailureCount": 102, "attributableFailureCount": 102,
@@ -6222,32 +6222,32 @@
"unresolvedFailureCount": 297 "unresolvedFailureCount": 297
}, },
"Sunrise_pt-200-x1": { "Sunrise_pt-200-x1": {
"attributableFailureCount": 732, "attributableFailureCount": 733,
"decisionFailureRate": 0.8894, "decisionFailureRate": 0.8896,
"decisionSuccessRate": 0.1106, "decisionSuccessRate": 0.1104,
"decisionTotal": 823, "decisionTotal": 824,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 162, "ambiguous_runtime": 162,
"architecture_compatibility": 44, "architecture_compatibility": 44,
"backend_operator": 35, "backend_operator": 36,
"context_length": 39, "context_length": 39,
"framework_architecture_unsupported": 154, "framework_architecture_unsupported": 154,
"memory_capacity": 113, "memory_capacity": 113,
"platform_infrastructure": 221, "platform_infrastructure": 222,
"repository_structure": 83, "repository_structure": 83,
"runtime_memory": 1, "runtime_memory": 1,
"tokenizer_compatibility": 263, "tokenizer_compatibility": 263,
"参数/模板问题": 228, "参数/模板问题": 228,
"验证失败": 32 "验证失败": 32
}, },
"failureCount": 1375, "failureCount": 1377,
"failureRate": 0.9379, "failureRate": 0.938,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 221, "platformFailureCount": 222,
"successCount": 91, "successCount": 91,
"successRate": 0.0621, "successRate": 0.062,
"total": 1466, "total": 1468,
"unresolvedFailureCount": 422 "unresolvedFailureCount": 422
}, },
"Vastai_va16": { "Vastai_va16": {
@@ -18908,14 +18908,14 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|awq": { "Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|awq": {
"attributableFailureCount": 5, "attributableFailureCount": 6,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 5, "decisionTotal": 6,
"failureBreakdown": { "failureBreakdown": {
"backend_operator": 5 "backend_operator": 6
}, },
"failureCount": 5, "failureCount": 6,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm_fix_tokenizer", "framework": "vllm_fix_tokenizer",
"modelType": "llama", "modelType": "llama",
@@ -18927,7 +18927,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1", "targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation", "taskType": "text-generation",
"total": 5, "total": 6,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|compressed-tensors": { "Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|compressed-tensors": {
@@ -21223,26 +21223,26 @@
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": { "Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"consecutiveFailures": 0, "consecutiveFailures": 0,
"consecutivePlatformFailures": 1, "consecutivePlatformFailures": 2,
"decisionFailureRate": 0.0, "decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 0, "decisionTotal": 0,
"failureBreakdown": { "failureBreakdown": {
"platform_infrastructure": 1 "platform_infrastructure": 2
}, },
"failureCount": 1, "failureCount": 2,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm_fix_tokenizer", "framework": "vllm_fix_tokenizer",
"lastPlatformFailureAt": "2026-09-23T01:50:08+00:00", "lastPlatformFailureAt": "2026-09-23T02:14:08+00:00",
"lastTerminalAt": "2026-09-23T01:50:30.356465+00:00", "lastTerminalAt": "2026-09-23T02:15:13.670721+00:00",
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 1, "platformFailureCount": 2,
"successCount": 0, "successCount": 0,
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1", "targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation", "taskType": "text-generation",
"total": 1, "total": 2,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"hygon_k100-ai|unknown|text-generation": { "hygon_k100-ai|unknown|text-generation": {
@@ -21625,31 +21625,6 @@
"total": 1, "total": 1,
"unresolvedFailureCount": 1 "unresolvedFailureCount": 1
}, },
"Biren_166m|vllm_fix_tokenizer|text-generation|gemma4|none": {
"attributableFailureCount": 0,
"consecutiveFailures": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"lastTerminalAt": "2026-09-21T12:58:50.741501+00:00",
"modelType": "gemma4",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Biren_166m",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"Biren_166m|vllm_fix_tokenizer|text-generation|hrm_text|none": { "Biren_166m|vllm_fix_tokenizer|text-generation|hrm_text|none": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"consecutiveFailures": 0, "consecutiveFailures": 0,
@@ -43262,14 +43237,14 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|awq|32": { "Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|awq|32": {
"attributableFailureCount": 4, "attributableFailureCount": 5,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 4, "decisionTotal": 5,
"failureBreakdown": { "failureBreakdown": {
"backend_operator": 4 "backend_operator": 5
}, },
"failureCount": 4, "failureCount": 5,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm_fix_tokenizer", "framework": "vllm_fix_tokenizer",
"loadSizeLog2Bucket": 32, "loadSizeLog2Bucket": 32,
@@ -43282,7 +43257,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1", "targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation", "taskType": "text-generation",
"total": 4, "total": 5,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|awq|33": { "Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|llama|awq|33": {
@@ -45755,23 +45730,23 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
} }
}, },
"terminalRecords": 16792, "terminalRecords": 16794,
"totalRecords": 17010, "totalRecords": 17012,
"totals": { "totals": {
"attributableFailureCount": 6011, "attributableFailureCount": 6012,
"decisionFailureRate": 0.8622, "decisionFailureRate": 0.8622,
"decisionSuccessRate": 0.1378, "decisionSuccessRate": 0.1378,
"decisionTotal": 6972, "decisionTotal": 6973,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 4282, "ambiguous_runtime": 4282,
"architecture_compatibility": 212, "architecture_compatibility": 212,
"attention_backend": 3, "attention_backend": 3,
"backend_operator": 113, "backend_operator": 114,
"context_length": 322, "context_length": 322,
"framework_architecture_unsupported": 2132, "framework_architecture_unsupported": 2132,
"memory_capacity": 1203, "memory_capacity": 1203,
"model_load": 526, "model_load": 526,
"platform_infrastructure": 930, "platform_infrastructure": 931,
"repository_structure": 740, "repository_structure": 740,
"runtime_memory": 91, "runtime_memory": 91,
"tokenizer_compatibility": 669, "tokenizer_compatibility": 669,
@@ -45779,20 +45754,19 @@
"日志缺失": 719, "日志缺失": 719,
"验证失败": 676 "验证失败": 676
}, },
"failureCount": 15831, "failureCount": 15833,
"failureRate": 0.9428, "failureRate": 0.9428,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 930, "platformFailureCount": 931,
"successCount": 961, "successCount": 961,
"successRate": 0.0572, "successRate": 0.0572,
"total": 16792, "total": 16794,
"unresolvedFailureCount": 8890 "unresolvedFailureCount": 8890
}, },
"warnings": [ "warnings": [
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。", "GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。", "GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。", "GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。", "GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
@@ -45801,8 +45775,9 @@
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。", "GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。", "GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。", "GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-100 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_bi-100 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Sunrise_pt-200-x1|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45826,7 +45801,6 @@
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45840,10 +45814,11 @@
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。" "组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
] ]
}, },
"storageMode": "decision_state_only", "storageMode": "decision_state_only",
"summarizedRecords": 17010, "summarizedRecords": 17012,
"version": 1 "version": 1
} }

View File

@@ -1,3 +1,4 @@
{"failReason": "platform_infrastructure", "failureAction": "open_short_gpu_framework_circuit_and_retry_other_models", "failureCategory": "platform_infrastructure", "failureClassificationReason": "no_idle_device", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-23T02:15:13.670721+00:00", "modelId": "ewinregirgojr/Qwen3.8-9B-Instruct-Turbo", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-23T02:14:08+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4516842", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "platform_infrastructure", "failureAction": "open_short_gpu_framework_circuit_and_retry_other_models", "failureCategory": "platform_infrastructure", "failureClassificationReason": "no_idle_device", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-23T01:50:30.356465+00:00", "modelId": "ewinregirgojr/Qwen3.8-14B-Instruct-Turbo", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-23T01:50:08+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4516843", "taskType": "text-generation", "verifyResult": -1} {"failReason": "platform_infrastructure", "failureAction": "open_short_gpu_framework_circuit_and_retry_other_models", "failureCategory": "platform_infrastructure", "failureClassificationReason": "no_idle_device", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-23T01:50:30.356465+00:00", "modelId": "ewinregirgojr/Qwen3.8-14B-Instruct-Turbo", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-23T01:50:08+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4516843", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-22T18:08:52.159386+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Quality-FP16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-22T18:06:08+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4592052", "taskType": "text-generation", "verifyResult": -1} {"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-22T18:08:52.159386+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Quality-FP16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-22T18:06:08+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4592052", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-22T10:15:00.768956+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-22T10:13:24+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4984023", "taskType": "text-generation", "verifyResult": null} {"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-22T10:15:00.768956+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-22T10:13:24+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4984023", "taskType": "text-generation", "verifyResult": null}
@@ -297,4 +298,3 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:25:05.959914+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293474987, "estimatedRequiredGiB": 18.232, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16313701503, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16313701503}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.969433+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986859", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:25:05.959914+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293474987, "estimatedRequiredGiB": 18.232, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16313701503, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16313701503}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.969433+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986859", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T12:09:51.363808+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.967582+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986857", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T12:09:51.363808+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.967582+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986857", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:21:25.760203+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683241105, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703487237, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703487237}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.966391+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986858", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:21:25.760203+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683241105, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703487237, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703487237}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.966391+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986858", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T12:58:50.741501+00:00", "modelId": "mlx-community/gemma-4-e2b-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5222073830, "estimatedRequiredGiB": 5.872, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 5254590634, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1140034851, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5254590634}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T19:48:14.804006+00:00", "targetGpu": "Biren_166m", "taskId": "4986623", "taskType": "text-generation", "verifyResult": -1}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -440,7 +440,6 @@
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-3.2-3B-Instruct-FP8", "modelId": "RedHatAI/Llama-3.2-3B-Instruct-FP8", "submitTime": "2026-09-20T09:25:37.077913+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976379", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-3.2-3B-Instruct-FP8", "modelId": "RedHatAI/Llama-3.2-3B-Instruct-FP8", "submitTime": "2026-09-20T09:25:37.077913+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976379", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "submitTime": "2026-09-20T09:25:37.090994+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976380", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "submitTime": "2026-09-20T09:25:37.090994+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976380", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/bharatgenai/Param2-17B-A2.4B-Thinking", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "submitTime": "2026-09-20T09:41:34.108703+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976623", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-ascend-910-b3"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/bharatgenai/Param2-17B-A2.4B-Thinking", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "submitTime": "2026-09-20T09:41:34.108703+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976623", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-ascend-910-b3"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-9b-it-quantized.w4a16", "modelId": "neuralmagic/gemma-2-9b-it-quantized.w4a16", "submitTime": "2026-09-20T10:49:41.634018+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4977490", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-kunlunxin-p-800"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-quantized.w8a8", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a8", "submitTime": "2026-09-20T11:19:51.535102+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4977812", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b4"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-quantized.w8a8", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a8", "submitTime": "2026-09-20T11:19:51.535102+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4977812", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b4"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "submitTime": "2026-09-20T11:19:51.641136+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4977829", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "submitTime": "2026-09-20T11:19:51.641136+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4977829", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "modelId": "neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "submitTime": "2026-09-20T11:52:06.534997+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4978249", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "modelId": "neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "submitTime": "2026-09-20T11:52:06.534997+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4978249", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
@@ -509,7 +508,6 @@
{"framework": "transformers", "modelAddress": "https://modelscope.cn/models/SyonLi/Qwen3-4B-Instruct-2507-Segmenter", "modelId": "SyonLi/Qwen3-4B-Instruct-2507-Segmenter", "submitTime": "2026-09-20T20:09:24.163788+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986867", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-transformers-iluvatar-bi-100"} {"framework": "transformers", "modelAddress": "https://modelscope.cn/models/SyonLi/Qwen3-4B-Instruct-2507-Segmenter", "modelId": "SyonLi/Qwen3-4B-Instruct-2507-Segmenter", "submitTime": "2026-09-20T20:09:24.163788+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986867", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-transformers-iluvatar-bi-100"}
{"framework": "transformers", "modelAddress": "https://modelscope.cn/models/mlx-community/LFM2.5-1.2B-Thinking-8bit", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-8bit", "submitTime": "2026-09-20T20:09:24.460163+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986875", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-transformers-iluvatar-bi-100"} {"framework": "transformers", "modelAddress": "https://modelscope.cn/models/mlx-community/LFM2.5-1.2B-Thinking-8bit", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-8bit", "submitTime": "2026-09-20T20:09:24.460163+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986875", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-transformers-iluvatar-bi-100"}
{"framework": "vllm-patch-tokenizer", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "submitTime": "2026-09-20T20:41:37.494440+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4987371", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-patch-tokenizer-hygon-k100-ai"} {"framework": "vllm-patch-tokenizer", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "submitTime": "2026-09-20T20:41:37.494440+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4987371", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-patch-tokenizer-hygon-k100-ai"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/AppleScript-8B-JANG_4M", "modelId": "JANGQ-AI/AppleScript-8B-JANG_4M", "submitTime": "2026-09-20T21:13:06.834956+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4987784", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-kunlunxin-p-800"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "submitTime": "2026-09-20T21:22:47.667525+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4987928", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-cambricon-mlu-370-x4"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "submitTime": "2026-09-20T21:22:47.667525+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4987928", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-cambricon-mlu-370-x4"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/aisingapore/Gemma-SEA-LION-v3-9B", "modelId": "aisingapore/Gemma-SEA-LION-v3-9B", "submitTime": "2026-09-20T21:26:40.552644+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4987963", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/aisingapore/Gemma-SEA-LION-v3-9B", "modelId": "aisingapore/Gemma-SEA-LION-v3-9B", "submitTime": "2026-09-20T21:26:40.552644+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4987963", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/mlx-community/Spark-X2.5-4B-OptiQ-4bit", "modelId": "mlx-community/Spark-X2.5-4B-OptiQ-4bit", "submitTime": "2026-09-20T21:28:22.564707+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4988004", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/mlx-community/Spark-X2.5-4B-OptiQ-4bit", "modelId": "mlx-community/Spark-X2.5-4B-OptiQ-4bit", "submitTime": "2026-09-20T21:28:22.564707+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4988004", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"}
@@ -557,7 +555,6 @@
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/BAAI/AquilaMed-RL", "modelId": "BAAI/AquilaMed-RL", "submitTime": "2026-09-21T04:44:11.759434+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4994057", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/BAAI/AquilaMed-RL", "modelId": "BAAI/AquilaMed-RL", "submitTime": "2026-09-21T04:44:11.759434+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4994057", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-9b-it-quantized.w4a16", "modelId": "RedHatAI/gemma-2-9b-it-quantized.w4a16", "submitTime": "2026-09-21T04:55:59.065763+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4994202", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-kunlunxin-p-800"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-9b-it-quantized.w4a16", "modelId": "RedHatAI/gemma-2-9b-it-quantized.w4a16", "submitTime": "2026-09-21T04:55:59.065763+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4994202", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-kunlunxin-p-800"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "submitTime": "2026-09-21T05:25:16.382748+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4994588", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-kunlunxin-p-800"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "submitTime": "2026-09-21T05:25:16.382748+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4994588", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-kunlunxin-p-800"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-e2b-it-qat-OptiQ-4bit", "modelId": "mlx-community/gemma-4-e2b-it-qat-OptiQ-4bit", "submitTime": "2026-09-21T05:25:16.370644+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4994586", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-kunlunxin-p-800"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/nm-testing/tinyllama-oneshot-w8w8-test-static-shape-change", "modelId": "nm-testing/tinyllama-oneshot-w8w8-test-static-shape-change", "submitTime": "2026-09-21T06:25:35.855289+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4995260", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b4"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/nm-testing/tinyllama-oneshot-w8w8-test-static-shape-change", "modelId": "nm-testing/tinyllama-oneshot-w8w8-test-static-shape-change", "submitTime": "2026-09-21T06:25:35.855289+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4995260", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b4"}
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w4a16", "modelId": "neuralmagic/gemma-2-2b-it-quantized.w4a16", "submitTime": "2026-09-21T06:49:13.744923+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4995536", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"} {"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w4a16", "modelId": "neuralmagic/gemma-2-2b-it-quantized.w4a16", "submitTime": "2026-09-21T06:49:13.744923+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4995536", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-3.2-3B-Instruct-FP8", "modelId": "RedHatAI/Llama-3.2-3B-Instruct-FP8", "submitTime": "2026-09-21T06:55:59.084023+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4995631", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"} {"framework": "vllm", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-3.2-3B-Instruct-FP8", "modelId": "RedHatAI/Llama-3.2-3B-Instruct-FP8", "submitTime": "2026-09-21T06:55:59.084023+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4995631", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"}

View File

@@ -2,24 +2,24 @@
"agentVersion": "2026.09.22.1", "agentVersion": "2026.09.22.1",
"checksums": { "checksums": {
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1", ".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
".modelhub_state/architecture_compatibility_blacklist.json": "81532c17d0752ac298a63a1cbc393e03077688c5267c17495eb1be9639922bbe", ".modelhub_state/architecture_compatibility_blacklist.json": "dff3c0eb8803cc8e84064e08d9b360b2ac0f4617d44a81289e00c6ff4c78bb52",
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab", ".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
".modelhub_state/market_intelligence.json": "f4b66e153e5c7663d2f2c3a3a32bbc3ca0d9984393f8d6cb0e3e97d711577d36", ".modelhub_state/market_intelligence.json": "9d957e9fa38d2ec780723cb770b360e5db0ecaf10083e03e135cc4a81f90e0cb",
".modelhub_state/official_capabilities.json": "19d6e7a9803a31704677a2b2e8541a283c561a9189ceebe09502305ea64bac70", ".modelhub_state/official_capabilities.json": "b34c80f0ea24c6f9389cb14b5bdfe6bbf1ae155af368dca418b9d9fbd736469a",
".modelhub_state/outcome_checkpoint.json": "94c4be1b7f953d240d4f8d664781d5fc2e1a598d37fdd86b02f2debef82a0f3c", ".modelhub_state/outcome_checkpoint.json": "4bf31db375a3fd79f0b3e75d3849476396251258109a129b3d611c40da4a6e1e",
".modelhub_state/queue_cleanup_latest.json": "abeedea8ce654c253250bcfd506d9b5eab392b53faf9867483b77b9a5c85233b", ".modelhub_state/queue_cleanup_latest.json": "abeedea8ce654c253250bcfd506d9b5eab392b53faf9867483b77b9a5c85233b",
".modelhub_state/recent_outcomes.jsonl": "6318aef4076b3471c99fbce38481ac745118c4f35da5bbd37517c3a2d621f343", ".modelhub_state/recent_outcomes.jsonl": "da9e62c7afd30c444050ea8df29de16a590f740faa5bfe281cda35a823b7283b",
".modelhub_state/recovery_active_tasks.jsonl": "22b4f1940e4315791f0244e9ee8a59dcffeb78ed662de9375e9e998ff506e3c2", ".modelhub_state/recovery_active_tasks.jsonl": "97b1d962075b044875e75700ec8b13e4695c844d1ed5e1c3825cae792b855583",
".modelhub_state/recovery_intents.jsonl": "cbad9c1acd40acb2a6943ea4872ce340b0fbfb5d217d4f570ff5333af0e680ff", ".modelhub_state/recovery_intents.jsonl": "7ab4a1b95bc12ed4178df4f187d87e54ed11cdaeecb731f01fe2ddacda9a062b",
".modelhub_state/routing_intelligence.json": "39ec7b5e77e11b96887ed828d68b094ab47d4d4e68b51e1ea92f4df7a81ab3e3", ".modelhub_state/routing_intelligence.json": "39ec7b5e77e11b96887ed828d68b094ab47d4d4e68b51e1ea92f4df7a81ab3e3",
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989", ".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
".modelhub_state/worker_crashes.jsonl": "619eac16f716b2420672872ff4ce92d452f5fced4f96e8dbb37f9ce511b53c5b", ".modelhub_state/worker_crashes.jsonl": "619eac16f716b2420672872ff4ce92d452f5fced4f96e8dbb37f9ce511b53c5b",
"ledger/submissions.jsonl": "53aef18e717c718994b0b46144b592a737928a181ffa9ef654e39cdbbb30185c", "ledger/submissions.jsonl": "c75a57c5121b2d8caff6e146ae9f83298543842c9f2490e1fd4eeb29a0738843",
"outcomes/submissions.jsonl": "5cc72ae4705f4e475770c78749c58b0e4a8f08739182dadcf2e2f2f00f38a1df" "outcomes/submissions.jsonl": "90f72d3ef137f3fc8dd279a856c1a6e746fcf32b32a548b3c315f9f9e135ae0d"
}, },
"generation": 13282, "generation": 13283,
"phase": "cycle", "phase": "cycle",
"schemaVersion": 1, "schemaVersion": 1,
"updatedAt": "2026-09-23T02:14:12.002514+00:00", "updatedAt": "2026-09-23T02:15:32.823905+00:00",
"writerId": "328f98096427428498653b568f0d5041" "writerId": "328f98096427428498653b568f0d5041"
} }

View File

@@ -2,7 +2,6 @@
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-04T12:29:56.904005+00:00", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "modelProfile": {"architectures": ["Param2MoEForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34302758832, "estimatedRequiredGiB": 38.353, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "param2moe", "modelscopeFileSize": 34317809830, "modelscopeLicense": null, "modelscopeParams": 17151125376, "modelscopeTags": ["model_type:param2moe", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:mixture-of-experts", "custom_tag:multilingual", "custom_tag:indian-languages"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 34317809830}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T04:29:16.672536+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4611191", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-04T12:29:56.904005+00:00", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "modelProfile": {"architectures": ["Param2MoEForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34302758832, "estimatedRequiredGiB": 38.353, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "param2moe", "modelscopeFileSize": 34317809830, "modelscopeLicense": null, "modelscopeParams": 17151125376, "modelscopeTags": ["model_type:param2moe", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:mixture-of-experts", "custom_tag:multilingual", "custom_tag:indian-languages"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 34317809830}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T04:29:16.672536+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4611191", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-04T12:43:05.302313+00:00", "modelId": "b77968543/Spark-X2.5-4B-Q8_0", "modelProfile": {"architectures": [], "configFingerprint": "054a7592edf89d789e18b12765c567c3c34b2ebb58661fdd49255aad03c66d40", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4375021152, "estimatedRequiredGiB": 4.889, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4375025003, "modelscopeLicense": null, "modelscopeParams": 4112079360, "modelscopeTags": ["library:gguf", "library:", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4375025003}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T04:40:16.372199+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4611339", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-04T12:43:05.302313+00:00", "modelId": "b77968543/Spark-X2.5-4B-Q8_0", "modelProfile": {"architectures": [], "configFingerprint": "054a7592edf89d789e18b12765c567c3c34b2ebb58661fdd49255aad03c66d40", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4375021152, "estimatedRequiredGiB": 4.889, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4375025003, "modelscopeLicense": null, "modelscopeParams": 4112079360, "modelscopeTags": ["library:gguf", "library:", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4375025003}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T04:40:16.372199+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4611339", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-04T22:53:17.234593+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 22125413673, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T14:46:47.882785+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4628625", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-04T22:53:17.234593+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 22125413673, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T14:46:47.882785+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4628625", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T00:55:08.690909+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T16:50:16.092995+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4630839", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T02:32:00.692571+00:00", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23420419700, "estimatedRequiredGiB": 26.214, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 23455515266, "modelscopeLicense": "mit", "modelscopeParams": 19528501104, "modelscopeTags": ["license:mit", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 23455515266}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T18:26:16.665549+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4632117", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T02:32:00.692571+00:00", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23420419700, "estimatedRequiredGiB": 26.214, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 23455515266, "modelscopeLicense": "mit", "modelscopeParams": 19528501104, "modelscopeTags": ["license:mit", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 23455515266}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T18:26:16.665549+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4632117", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T03:05:43.096396+00:00", "modelId": "neuralmagic/gemma-2-9b-it-quantized.w8a8", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 11998609696, "estimatedRequiredGiB": 13.434, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 12020499228, "modelscopeLicense": "gemma", "modelscopeParams": 10159209984, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 12020499228}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T19:03:47.678567+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4632579", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T03:05:43.096396+00:00", "modelId": "neuralmagic/gemma-2-9b-it-quantized.w8a8", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 11998609696, "estimatedRequiredGiB": 13.434, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 12020499228, "modelscopeLicense": "gemma", "modelscopeParams": 10159209984, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 12020499228}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T19:03:47.678567+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4632579", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T03:32:15.596901+00:00", "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4395007696, "estimatedRequiredGiB": 4.922, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 4404161088, "modelscopeLicense": "llama3.2", "modelscopeParams": 3606752256, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4404161088}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T19:29:42.987166+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4632974", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T03:32:15.596901+00:00", "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4395007696, "estimatedRequiredGiB": 4.922, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 4404161088, "modelscopeLicense": "llama3.2", "modelscopeParams": 3606752256, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4404161088}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T19:29:42.987166+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4632974", "taskType": "text-generation", "verifyResult": null}