state: generation 13298 (cycle)

This commit is contained in:
2026-09-23 02:35:51 +00:00
parent ca26bbc89e
commit 4b3af6c8ac
9 changed files with 1719 additions and 1696 deletions

View File

@@ -3044,7 +3044,7 @@
"taskType": "text-generation" "taskType": "text-generation"
} }
}, },
"generatedAt": "2026-09-23T02:31:29.971650+00:00", "generatedAt": "2026-09-23T02:35:36.844851+00:00",
"summary": { "summary": {
"activeBlockCount": 153, "activeBlockCount": 153,
"byGpuFramework": { "byGpuFramework": {

View File

@@ -396,7 +396,7 @@
} }
}, },
"frameworkUpdatedAt": null, "frameworkUpdatedAt": null,
"generatedAt": "2026-09-23T02:34:31.553646+00:00", "generatedAt": "2026-09-23T02:35:49.709082+00:00",
"gpuStats": { "gpuStats": {
"Ascend_910-b4": { "Ascend_910-b4": {
"available": true, "available": true,

View File

@@ -1,5 +1,5 @@
{ {
"catalogUpdatedAt": "2026-09-23T02:34:31.553646+00:00", "catalogUpdatedAt": "2026-09-23T02:35:49.709082+00:00",
"configuredTaskTypes": [ "configuredTaskTypes": [
"text-generation" "text-generation"
], ],
@@ -56,7 +56,7 @@
"time-series-forecasting" "time-series-forecasting"
], ],
"errors": [], "errors": [],
"generatedAt": "2026-09-23T02:34:33.860576+00:00", "generatedAt": "2026-09-23T02:35:49.709082+00:00",
"gpuCatalog": { "gpuCatalog": {
"Ascend_910-b3": { "Ascend_910-b3": {
"canVerify": true, "canVerify": true,
@@ -6693,6 +6693,6 @@
"updateTime": "2025-12-22 08:59:53" "updateTime": "2025-12-22 08:59:53"
} }
], ],
"taskTreeUpdatedAt": "2026-09-23T02:34:31.553646+00:00", "taskTreeUpdatedAt": "2026-09-23T02:35:49.709082+00:00",
"version": 1 "version": 1
} }

View File

@@ -1,6 +1,6 @@
{ {
"generatedAt": "2026-09-23T02:31:29.894121+00:00", "generatedAt": "2026-09-23T02:35:36.757931+00:00",
"lastSyncTime": "2026-09-23T02:31:28.170137+00:00", "lastSyncTime": "2026-09-23T02:35:34.755698+00:00",
"recentLimit": 300, "recentLimit": 300,
"report": { "report": {
"architectureCompatibilityBlocks": { "architectureCompatibilityBlocks": {
@@ -4729,17 +4729,17 @@
"unresolvedFailureCount": 19 "unresolvedFailureCount": 19
}, },
"Kunlunxin_p-800|vllm|text-generation": { "Kunlunxin_p-800|vllm|text-generation": {
"attributableFailureCount": 12, "attributableFailureCount": 13,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 12, "decisionTotal": 13,
"failureBreakdown": { "failureBreakdown": {
"backend_operator": 2, "backend_operator": 3,
"framework_architecture_unsupported": 9, "framework_architecture_unsupported": 9,
"model_load": 1, "model_load": 1,
"参数/模板问题": 2 "参数/模板问题": 2
}, },
"failureCount": 14, "failureCount": 15,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm", "framework": "vllm",
"pendingCount": 0, "pendingCount": 0,
@@ -4749,7 +4749,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Kunlunxin_p-800", "targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation", "taskType": "text-generation",
"total": 14, "total": 15,
"unresolvedFailureCount": 2 "unresolvedFailureCount": 2
}, },
"Kunlunxin_r-200-8f|unknown|text-generation": { "Kunlunxin_r-200-8f|unknown|text-generation": {
@@ -5712,15 +5712,15 @@
"unresolvedFailureCount": 6462 "unresolvedFailureCount": 6462
}, },
"vllm": { "vllm": {
"attributableFailureCount": 3574, "attributableFailureCount": 3575,
"decisionFailureRate": 0.9741, "decisionFailureRate": 0.9741,
"decisionSuccessRate": 0.0259, "decisionSuccessRate": 0.0259,
"decisionTotal": 3669, "decisionTotal": 3670,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 1534, "ambiguous_runtime": 1534,
"architecture_compatibility": 112, "architecture_compatibility": 112,
"attention_backend": 3, "attention_backend": 3,
"backend_operator": 86, "backend_operator": 87,
"context_length": 161, "context_length": 161,
"framework_architecture_unsupported": 1330, "framework_architecture_unsupported": 1330,
"memory_capacity": 722, "memory_capacity": 722,
@@ -5731,14 +5731,14 @@
"tokenizer_compatibility": 413, "tokenizer_compatibility": 413,
"参数/模板问题": 48 "参数/模板问题": 48
}, },
"failureCount": 6022, "failureCount": 6023,
"failureRate": 0.9845, "failureRate": 0.9845,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 866, "platformFailureCount": 866,
"successCount": 95, "successCount": 95,
"successRate": 0.0155, "successRate": 0.0155,
"total": 6117, "total": 6118,
"unresolvedFailureCount": 1582 "unresolvedFailureCount": 1582
}, },
"vllm-customized": { "vllm-customized": {
@@ -5888,7 +5888,7 @@
"unresolvedFailureCount": 175 "unresolvedFailureCount": 175
} }
}, },
"generatedAt": "2026-09-23T02:31:29.882498+00:00", "generatedAt": "2026-09-23T02:35:36.744395+00:00",
"gpuSummaries": { "gpuSummaries": {
"Ascend_910-b3": { "Ascend_910-b3": {
"attributableFailureCount": 102, "attributableFailureCount": 102,
@@ -6114,13 +6114,13 @@
"unresolvedFailureCount": 717 "unresolvedFailureCount": 717
}, },
"Kunlunxin_p-800": { "Kunlunxin_p-800": {
"attributableFailureCount": 53, "attributableFailureCount": 54,
"decisionFailureRate": 0.9138, "decisionFailureRate": 0.9153,
"decisionSuccessRate": 0.0862, "decisionSuccessRate": 0.0847,
"decisionTotal": 58, "decisionTotal": 59,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 210, "ambiguous_runtime": 210,
"backend_operator": 10, "backend_operator": 11,
"framework_architecture_unsupported": 24, "framework_architecture_unsupported": 24,
"memory_capacity": 2, "memory_capacity": 2,
"model_load": 13, "model_load": 13,
@@ -6131,14 +6131,14 @@
"日志缺失": 1, "日志缺失": 1,
"验证失败": 23 "验证失败": 23
}, },
"failureCount": 320, "failureCount": 321,
"failureRate": 0.9846, "failureRate": 0.9847,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 1, "platformFailureCount": 1,
"successCount": 5, "successCount": 5,
"successRate": 0.0154, "successRate": 0.0153,
"total": 325, "total": 326,
"unresolvedFailureCount": 266 "unresolvedFailureCount": 266
}, },
"Kunlunxin_r-200-8f": { "Kunlunxin_r-200-8f": {
@@ -17678,14 +17678,14 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": { "Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": {
"attributableFailureCount": 2, "attributableFailureCount": 3,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 2, "decisionTotal": 3,
"failureBreakdown": { "failureBreakdown": {
"backend_operator": 2 "backend_operator": 3
}, },
"failureCount": 2, "failureCount": 3,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm", "framework": "vllm",
"modelType": "starcoder2", "modelType": "starcoder2",
@@ -17697,7 +17697,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Kunlunxin_p-800", "targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation", "taskType": "text-generation",
"total": 2, "total": 3,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"MetaX_c-500|vllm|text-generation|aquila3|none": { "MetaX_c-500|vllm|text-generation|aquila3|none": {
@@ -20888,9 +20888,9 @@
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 0, "decisionTotal": 0,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 16 "ambiguous_runtime": 15
}, },
"failureCount": 16, "failureCount": 15,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "transformers", "framework": "transformers",
"lastPlatformFailureAt": null, "lastPlatformFailureAt": null,
@@ -20902,8 +20902,8 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Iluvatar_bi-100", "targetGpu": "Iluvatar_bi-100",
"taskType": "text-generation", "taskType": "text-generation",
"total": 16, "total": 15,
"unresolvedFailureCount": 16 "unresolvedFailureCount": 15
}, },
"Iluvatar_bi-150|transformers|text-generation": { "Iluvatar_bi-150|transformers|text-generation": {
"attributableFailureCount": 5, "attributableFailureCount": 5,
@@ -21141,18 +21141,18 @@
"unresolvedFailureCount": 1 "unresolvedFailureCount": 1
}, },
"Kunlunxin_p-800|vllm|text-generation": { "Kunlunxin_p-800|vllm|text-generation": {
"attributableFailureCount": 6, "attributableFailureCount": 7,
"consecutiveFailures": 6, "consecutiveFailures": 7,
"consecutivePlatformFailures": 0, "consecutivePlatformFailures": 0,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 6, "decisionTotal": 7,
"failureBreakdown": { "failureBreakdown": {
"backend_operator": 2, "backend_operator": 3,
"framework_architecture_unsupported": 4, "framework_architecture_unsupported": 4,
"参数/模板问题": 2 "参数/模板问题": 2
}, },
"failureCount": 8, "failureCount": 9,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm", "framework": "vllm",
"lastPlatformFailureAt": null, "lastPlatformFailureAt": null,
@@ -21164,7 +21164,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Kunlunxin_p-800", "targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation", "taskType": "text-generation",
"total": 8, "total": 9,
"unresolvedFailureCount": 2 "unresolvedFailureCount": 2
}, },
"MetaX_c-500|vllm|text-generation": { "MetaX_c-500|vllm|text-generation": {
@@ -23005,9 +23005,9 @@
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 0, "decisionTotal": 0,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 2 "ambiguous_runtime": 1
}, },
"failureCount": 2, "failureCount": 1,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "transformers", "framework": "transformers",
"lastTerminalAt": "2026-09-21T05:42:03.443562+00:00", "lastTerminalAt": "2026-09-21T05:42:03.443562+00:00",
@@ -23020,8 +23020,8 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Iluvatar_bi-100", "targetGpu": "Iluvatar_bi-100",
"taskType": "text-generation", "taskType": "text-generation",
"total": 2, "total": 1,
"unresolvedFailureCount": 2 "unresolvedFailureCount": 1
}, },
"Iluvatar_bi-100|transformers|text-generation|mistral3|none": { "Iluvatar_bi-100|transformers|text-generation|mistral3|none": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
@@ -24150,18 +24150,18 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": { "Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": {
"attributableFailureCount": 2, "attributableFailureCount": 3,
"consecutiveFailures": 2, "consecutiveFailures": 3,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 2, "decisionTotal": 3,
"failureBreakdown": { "failureBreakdown": {
"backend_operator": 2 "backend_operator": 3
}, },
"failureCount": 2, "failureCount": 3,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm", "framework": "vllm",
"lastTerminalAt": "2026-09-22T08:25:19.951625+00:00", "lastTerminalAt": "2026-09-23T02:35:34.755698+00:00",
"modelType": "starcoder2", "modelType": "starcoder2",
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
@@ -24171,7 +24171,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "Kunlunxin_p-800", "targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation", "taskType": "text-generation",
"total": 2, "total": 3,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Mthreads_s4000|vllm|text-generation|llama|compressed-tensors": { "Mthreads_s4000|vllm|text-generation|llama|compressed-tensors": {
@@ -41372,6 +41372,30 @@
"total": 2, "total": 2,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors|34": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"backend_operator": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"loadSizeLog2Bucket": 34,
"modelType": "starcoder2",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"MetaX_c-500|vllm|text-generation|aquila3|none|12": { "MetaX_c-500|vllm|text-generation|aquila3|none|12": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"decisionFailureRate": 0.0, "decisionFailureRate": 0.0,
@@ -45802,18 +45826,18 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
} }
}, },
"terminalRecords": 16795, "terminalRecords": 16796,
"totalRecords": 17013, "totalRecords": 17014,
"totals": { "totals": {
"attributableFailureCount": 6012, "attributableFailureCount": 6013,
"decisionFailureRate": 0.8622, "decisionFailureRate": 0.8622,
"decisionSuccessRate": 0.1378, "decisionSuccessRate": 0.1378,
"decisionTotal": 6973, "decisionTotal": 6974,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 4283, "ambiguous_runtime": 4283,
"architecture_compatibility": 212, "architecture_compatibility": 212,
"attention_backend": 3, "attention_backend": 3,
"backend_operator": 114, "backend_operator": 115,
"context_length": 322, "context_length": 322,
"framework_architecture_unsupported": 2132, "framework_architecture_unsupported": 2132,
"memory_capacity": 1203, "memory_capacity": 1203,
@@ -45826,19 +45850,20 @@
"日志缺失": 719, "日志缺失": 719,
"验证失败": 676 "验证失败": 676
}, },
"failureCount": 15834, "failureCount": 15835,
"failureRate": 0.9428, "failureRate": 0.9428,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 931, "platformFailureCount": 931,
"successCount": 961, "successCount": 961,
"successRate": 0.0572, "successRate": 0.0572,
"total": 16795, "total": 16796,
"unresolvedFailureCount": 8891 "unresolvedFailureCount": 8891
}, },
"warnings": [ "warnings": [
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。", "GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。", "GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。", "GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。", "GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
@@ -45849,7 +45874,6 @@
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。", "GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。", "GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。", "GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-100 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_bi-100 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Sunrise_pt-200-x1|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45873,7 +45897,6 @@
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45887,10 +45910,11 @@
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。" "组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
] ]
}, },
"storageMode": "decision_state_only", "storageMode": "decision_state_only",
"summarizedRecords": 17013, "summarizedRecords": 17014,
"version": 1 "version": 1
} }

View File

@@ -76,6 +76,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T04:22:26.358545+00:00", "modelId": "webAI-Official/TwIL-LM3", "modelProfile": {"architectures": ["SmolLM3ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6150235096, "estimatedRequiredGiB": 24.879, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "smollm3", "modelscopeFileSize": 22261451381, "modelscopeLicense": "other", "modelscopeParams": 3075098624, "modelscopeTags": ["license:other", "model_type:smollm3", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:smollm3", "custom_tag:twil-lm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22261451381}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:17:59.369215+00:00", "targetGpu": "Biren_166m", "taskId": "5002150", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T04:22:26.358545+00:00", "modelId": "webAI-Official/TwIL-LM3", "modelProfile": {"architectures": ["SmolLM3ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6150235096, "estimatedRequiredGiB": 24.879, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "smollm3", "modelscopeFileSize": 22261451381, "modelscopeLicense": "other", "modelscopeParams": 3075098624, "modelscopeTags": ["license:other", "model_type:smollm3", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:smollm3", "custom_tag:twil-lm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22261451381}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:17:59.369215+00:00", "targetGpu": "Biren_166m", "taskId": "5002150", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T16:20:26.367867+00:00", "modelId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-8bit", "modelProfile": {"architectures": ["MuseGlimmerForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 35620174949, "estimatedRequiredGiB": 39.84, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "muse_glimmer", "modelscopeFileSize": 35648665467, "modelscopeLicense": "apache-2.0", "modelscopeParams": 9345525760, "modelscopeTags": ["license:apache-2.0", "model_type:muse_glimmer", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:muse_glimmer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 35648665467}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T13:43:30.755806+00:00", "targetGpu": "Biren_166m", "taskId": "5001736", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T16:20:26.367867+00:00", "modelId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-8bit", "modelProfile": {"architectures": ["MuseGlimmerForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 35620174949, "estimatedRequiredGiB": 39.84, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "muse_glimmer", "modelscopeFileSize": 35648665467, "modelscopeLicense": "apache-2.0", "modelscopeParams": 9345525760, "modelscopeTags": ["license:apache-2.0", "model_type:muse_glimmer", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:muse_glimmer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 35648665467}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T13:43:30.755806+00:00", "targetGpu": "Biren_166m", "taskId": "5001736", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-23T00:07:02.456914+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:49:01.244975+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001063", "taskType": "text-generation", "verifyResult": null} {"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-23T00:07:02.456914+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:49:01.244975+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001063", "taskType": "text-generation", "verifyResult": null}
{"failReason": "backend_operator", "failureAction": "prefer_other_proven_framework_or_gpu", "failureCategory": "backend_operator", "failureClassificationReason": "backend_version_dependent", "failureCode": "MISSING_OPERATOR", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-23T02:35:34.755698+00:00", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663867, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663867}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:49:01.235875+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001055", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "参数/模板问题", "framework": "vllm-customized", "lastSyncTime": "2026-09-22T09:59:25.057632+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:25:50.997483+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000706", "taskType": "text-generation", "verifyResult": null} {"failReason": "参数/模板问题", "framework": "vllm-customized", "lastSyncTime": "2026-09-22T09:59:25.057632+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:25:50.997483+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000706", "taskType": "text-generation", "verifyResult": null}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859282+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.453427+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000553", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859282+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.453427+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000553", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859263+00:00", "modelId": "mlx-community/gemma-4-e2b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5226595850, "estimatedRequiredGiB": 5.878, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 5259113427, "modelscopeLicense": "gemma", "modelscopeParams": 1141165347, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5259113427}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.451262+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000554", "taskType": "text-generation", "verifyResult": -1} {"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859263+00:00", "modelId": "mlx-community/gemma-4-e2b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5226595850, "estimatedRequiredGiB": 5.878, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 5259113427, "modelscopeLicense": "gemma", "modelscopeParams": 1141165347, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5259113427}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.451262+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000554", "taskType": "text-generation", "verifyResult": -1}
@@ -297,4 +298,3 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457647+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.994993+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986861", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457647+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.994993+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986861", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_vl"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:14:26.178418+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.991772+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986860", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_vl"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:14:26.178418+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.991772+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986860", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:25:05.959914+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293474987, "estimatedRequiredGiB": 18.232, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16313701503, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16313701503}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.969433+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986859", "taskType": "text-generation", "verifyResult": -1} {"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:25:05.959914+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293474987, "estimatedRequiredGiB": 18.232, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16313701503, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16313701503}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.969433+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986859", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T12:09:51.363808+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.967582+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986857", "taskType": "text-generation", "verifyResult": -1}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -2,24 +2,24 @@
"agentVersion": "2026.09.22.1", "agentVersion": "2026.09.22.1",
"checksums": { "checksums": {
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1", ".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
".modelhub_state/architecture_compatibility_blacklist.json": "cdcfe698feaf16eac366b1264402fa79c5d03a36f254171e091485de4f782fc1", ".modelhub_state/architecture_compatibility_blacklist.json": "d891b45df03e4a1e56761b86b5162b5d74120bec1c693327f9b2785286a71088",
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab", ".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
".modelhub_state/market_intelligence.json": "fd79be793623b5d9f3968d56b476a9f9956ddd811ee63924f6e96f354ee72db0", ".modelhub_state/market_intelligence.json": "55813e4d28cd521b087a47128ba7d1842d31418b407e8f8c8a237fc016933f6b",
".modelhub_state/official_capabilities.json": "2ace5dc16013b6195381aad9051bdae71e5eb59ae38a8486575b22d2dbd95321", ".modelhub_state/official_capabilities.json": "602abbdd56c00692f3425c589e14fd39889c418bf3e08b5e63f3cb5d33439e57",
".modelhub_state/outcome_checkpoint.json": "207ee0fe96f1ad0125e5ee17042986e25d246e41e593ec45452569175a757808", ".modelhub_state/outcome_checkpoint.json": "eb01f90f8b326f3cb5ea18888009e1ae273920b6ab85f2331e475a61609a3223",
".modelhub_state/queue_cleanup_latest.json": "abeedea8ce654c253250bcfd506d9b5eab392b53faf9867483b77b9a5c85233b", ".modelhub_state/queue_cleanup_latest.json": "abeedea8ce654c253250bcfd506d9b5eab392b53faf9867483b77b9a5c85233b",
".modelhub_state/recent_outcomes.jsonl": "9dfc138ee2bbd5fda6efa8088e67b8fa5b977fccc750b1c9b3a524cce64d9fea", ".modelhub_state/recent_outcomes.jsonl": "be89f770ae4bdc0e9de6f4ee70ec569bf6dea08f70738c7d578a4ebea2bf310a",
".modelhub_state/recovery_active_tasks.jsonl": "9e82b697561031e34efa522cf1d8b38506deb81594f81fb548e096631229776b", ".modelhub_state/recovery_active_tasks.jsonl": "10159aaebad20748935c0aabde549201d4b154b6e1fc11476f67418862a9ce1c",
".modelhub_state/recovery_intents.jsonl": "66ef539db5857f9619c99c01c602096c0e3cc18d1fee3998238e63e06aefa88c", ".modelhub_state/recovery_intents.jsonl": "ddd029a6dbae152d255f2bda464b2c17d3523e15f7f6cca2ffc394735a45dc38",
".modelhub_state/routing_intelligence.json": "39ec7b5e77e11b96887ed828d68b094ab47d4d4e68b51e1ea92f4df7a81ab3e3", ".modelhub_state/routing_intelligence.json": "39ec7b5e77e11b96887ed828d68b094ab47d4d4e68b51e1ea92f4df7a81ab3e3",
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989", ".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
".modelhub_state/worker_crashes.jsonl": "619eac16f716b2420672872ff4ce92d452f5fced4f96e8dbb37f9ce511b53c5b", ".modelhub_state/worker_crashes.jsonl": "619eac16f716b2420672872ff4ce92d452f5fced4f96e8dbb37f9ce511b53c5b",
"ledger/submissions.jsonl": "1949681660efa39796ca1a4f1338205cd3a2d7ae211571a2b78e7fc663fcba49", "ledger/submissions.jsonl": "1949681660efa39796ca1a4f1338205cd3a2d7ae211571a2b78e7fc663fcba49",
"outcomes/submissions.jsonl": "0f7921d4c479dda8e98cb0985f6b6d1b3b1615de93ab424b1357bac650a23583" "outcomes/submissions.jsonl": "860e986a50ebcb1a9c0bd01cf90dca36eeb907a901d39e16e56191a0f3f38c0e"
}, },
"generation": 13297, "generation": 13298,
"phase": "cycle", "phase": "cycle",
"schemaVersion": 1, "schemaVersion": 1,
"updatedAt": "2026-09-23T02:34:34.092787+00:00", "updatedAt": "2026-09-23T02:35:51.180630+00:00",
"writerId": "328f98096427428498653b568f0d5041" "writerId": "328f98096427428498653b568f0d5041"
} }

View File

@@ -678,7 +678,6 @@
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063157+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020365960, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022810352, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022810352}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.259499+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001065", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063157+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020365960, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022810352, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022810352}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.259499+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001065", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063192+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020349648, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022916227, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022916227}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.241568+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001058", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063192+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020349648, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022916227, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022916227}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.241568+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001058", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063175+00:00", "modelId": "BAAI/AquilaMed-RL", "modelProfile": {"architectures": ["AquilaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6968, "estimatedRequiredGiB": 18.386, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "aquila3", "modelscopeFileSize": 16451477782, "modelscopeLicense": "other", "modelscopeParams": 8223748096, "modelscopeTags": ["license:other", "model_type:aquila3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:arxiv:2406.12182"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16451477782}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.257127+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001059", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063175+00:00", "modelId": "BAAI/AquilaMed-RL", "modelProfile": {"architectures": ["AquilaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6968, "estimatedRequiredGiB": 18.386, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "aquila3", "modelscopeFileSize": 16451477782, "modelscopeLicense": "other", "modelscopeParams": 8223748096, "modelscopeTags": ["license:other", "model_type:aquila3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:arxiv:2406.12182"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16451477782}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.257127+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001059", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063143+00:00", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663867, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663867}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.235875+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001055", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063105+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B", "modelProfile": {"architectures": ["Lfm2MoeForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16936006912, "estimatedRequiredGiB": 18.948, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2_moe", "modelscopeFileSize": 16953948786, "modelscopeLicense": "other", "modelscopeParams": 8467856832, "modelscopeTags": ["license:other", "model_type:lfm2_moe", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16953948786}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.246640+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001066", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063105+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B", "modelProfile": {"architectures": ["Lfm2MoeForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16936006912, "estimatedRequiredGiB": 18.948, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2_moe", "modelscopeFileSize": 16953948786, "modelscopeLicense": "other", "modelscopeParams": 8467856832, "modelscopeTags": ["license:other", "model_type:lfm2_moe", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16953948786}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.246640+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001066", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063182+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.237142+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001062", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063182+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.237142+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001062", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063150+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 42709318472, "estimatedRequiredGiB": 47.753, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 42728540075, "modelscopeLicense": "llama3.1", "modelscopeParams": 70553707056, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 42728540075}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.249214+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001061", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063150+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 42709318472, "estimatedRequiredGiB": 47.753, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 42728540075, "modelscopeLicense": "llama3.1", "modelscopeParams": 70553707056, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 42728540075}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.249214+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001061", "taskType": "text-generation", "verifyResult": null}