state: generation 13298 (cycle)
This commit is contained in:
@@ -3044,7 +3044,7 @@
|
|||||||
"taskType": "text-generation"
|
"taskType": "text-generation"
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"generatedAt": "2026-09-23T02:31:29.971650+00:00",
|
"generatedAt": "2026-09-23T02:35:36.844851+00:00",
|
||||||
"summary": {
|
"summary": {
|
||||||
"activeBlockCount": 153,
|
"activeBlockCount": 153,
|
||||||
"byGpuFramework": {
|
"byGpuFramework": {
|
||||||
|
|||||||
@@ -396,7 +396,7 @@
|
|||||||
}
|
}
|
||||||
},
|
},
|
||||||
"frameworkUpdatedAt": null,
|
"frameworkUpdatedAt": null,
|
||||||
"generatedAt": "2026-09-23T02:34:31.553646+00:00",
|
"generatedAt": "2026-09-23T02:35:49.709082+00:00",
|
||||||
"gpuStats": {
|
"gpuStats": {
|
||||||
"Ascend_910-b4": {
|
"Ascend_910-b4": {
|
||||||
"available": true,
|
"available": true,
|
||||||
|
|||||||
@@ -1,5 +1,5 @@
|
|||||||
{
|
{
|
||||||
"catalogUpdatedAt": "2026-09-23T02:34:31.553646+00:00",
|
"catalogUpdatedAt": "2026-09-23T02:35:49.709082+00:00",
|
||||||
"configuredTaskTypes": [
|
"configuredTaskTypes": [
|
||||||
"text-generation"
|
"text-generation"
|
||||||
],
|
],
|
||||||
@@ -56,7 +56,7 @@
|
|||||||
"time-series-forecasting"
|
"time-series-forecasting"
|
||||||
],
|
],
|
||||||
"errors": [],
|
"errors": [],
|
||||||
"generatedAt": "2026-09-23T02:34:33.860576+00:00",
|
"generatedAt": "2026-09-23T02:35:49.709082+00:00",
|
||||||
"gpuCatalog": {
|
"gpuCatalog": {
|
||||||
"Ascend_910-b3": {
|
"Ascend_910-b3": {
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
@@ -6693,6 +6693,6 @@
|
|||||||
"updateTime": "2025-12-22 08:59:53"
|
"updateTime": "2025-12-22 08:59:53"
|
||||||
}
|
}
|
||||||
],
|
],
|
||||||
"taskTreeUpdatedAt": "2026-09-23T02:34:31.553646+00:00",
|
"taskTreeUpdatedAt": "2026-09-23T02:35:49.709082+00:00",
|
||||||
"version": 1
|
"version": 1
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -1,6 +1,6 @@
|
|||||||
{
|
{
|
||||||
"generatedAt": "2026-09-23T02:31:29.894121+00:00",
|
"generatedAt": "2026-09-23T02:35:36.757931+00:00",
|
||||||
"lastSyncTime": "2026-09-23T02:31:28.170137+00:00",
|
"lastSyncTime": "2026-09-23T02:35:34.755698+00:00",
|
||||||
"recentLimit": 300,
|
"recentLimit": 300,
|
||||||
"report": {
|
"report": {
|
||||||
"architectureCompatibilityBlocks": {
|
"architectureCompatibilityBlocks": {
|
||||||
@@ -4729,17 +4729,17 @@
|
|||||||
"unresolvedFailureCount": 19
|
"unresolvedFailureCount": 19
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800|vllm|text-generation": {
|
"Kunlunxin_p-800|vllm|text-generation": {
|
||||||
"attributableFailureCount": 12,
|
"attributableFailureCount": 13,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 12,
|
"decisionTotal": 13,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"backend_operator": 2,
|
"backend_operator": 3,
|
||||||
"framework_architecture_unsupported": 9,
|
"framework_architecture_unsupported": 9,
|
||||||
"model_load": 1,
|
"model_load": 1,
|
||||||
"参数/模板问题": 2
|
"参数/模板问题": 2
|
||||||
},
|
},
|
||||||
"failureCount": 14,
|
"failureCount": 15,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
@@ -4749,7 +4749,7 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Kunlunxin_p-800",
|
"targetGpu": "Kunlunxin_p-800",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 14,
|
"total": 15,
|
||||||
"unresolvedFailureCount": 2
|
"unresolvedFailureCount": 2
|
||||||
},
|
},
|
||||||
"Kunlunxin_r-200-8f|unknown|text-generation": {
|
"Kunlunxin_r-200-8f|unknown|text-generation": {
|
||||||
@@ -5712,15 +5712,15 @@
|
|||||||
"unresolvedFailureCount": 6462
|
"unresolvedFailureCount": 6462
|
||||||
},
|
},
|
||||||
"vllm": {
|
"vllm": {
|
||||||
"attributableFailureCount": 3574,
|
"attributableFailureCount": 3575,
|
||||||
"decisionFailureRate": 0.9741,
|
"decisionFailureRate": 0.9741,
|
||||||
"decisionSuccessRate": 0.0259,
|
"decisionSuccessRate": 0.0259,
|
||||||
"decisionTotal": 3669,
|
"decisionTotal": 3670,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 1534,
|
"ambiguous_runtime": 1534,
|
||||||
"architecture_compatibility": 112,
|
"architecture_compatibility": 112,
|
||||||
"attention_backend": 3,
|
"attention_backend": 3,
|
||||||
"backend_operator": 86,
|
"backend_operator": 87,
|
||||||
"context_length": 161,
|
"context_length": 161,
|
||||||
"framework_architecture_unsupported": 1330,
|
"framework_architecture_unsupported": 1330,
|
||||||
"memory_capacity": 722,
|
"memory_capacity": 722,
|
||||||
@@ -5731,14 +5731,14 @@
|
|||||||
"tokenizer_compatibility": 413,
|
"tokenizer_compatibility": 413,
|
||||||
"参数/模板问题": 48
|
"参数/模板问题": 48
|
||||||
},
|
},
|
||||||
"failureCount": 6022,
|
"failureCount": 6023,
|
||||||
"failureRate": 0.9845,
|
"failureRate": 0.9845,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 866,
|
"platformFailureCount": 866,
|
||||||
"successCount": 95,
|
"successCount": 95,
|
||||||
"successRate": 0.0155,
|
"successRate": 0.0155,
|
||||||
"total": 6117,
|
"total": 6118,
|
||||||
"unresolvedFailureCount": 1582
|
"unresolvedFailureCount": 1582
|
||||||
},
|
},
|
||||||
"vllm-customized": {
|
"vllm-customized": {
|
||||||
@@ -5888,7 +5888,7 @@
|
|||||||
"unresolvedFailureCount": 175
|
"unresolvedFailureCount": 175
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"generatedAt": "2026-09-23T02:31:29.882498+00:00",
|
"generatedAt": "2026-09-23T02:35:36.744395+00:00",
|
||||||
"gpuSummaries": {
|
"gpuSummaries": {
|
||||||
"Ascend_910-b3": {
|
"Ascend_910-b3": {
|
||||||
"attributableFailureCount": 102,
|
"attributableFailureCount": 102,
|
||||||
@@ -6114,13 +6114,13 @@
|
|||||||
"unresolvedFailureCount": 717
|
"unresolvedFailureCount": 717
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800": {
|
"Kunlunxin_p-800": {
|
||||||
"attributableFailureCount": 53,
|
"attributableFailureCount": 54,
|
||||||
"decisionFailureRate": 0.9138,
|
"decisionFailureRate": 0.9153,
|
||||||
"decisionSuccessRate": 0.0862,
|
"decisionSuccessRate": 0.0847,
|
||||||
"decisionTotal": 58,
|
"decisionTotal": 59,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 210,
|
"ambiguous_runtime": 210,
|
||||||
"backend_operator": 10,
|
"backend_operator": 11,
|
||||||
"framework_architecture_unsupported": 24,
|
"framework_architecture_unsupported": 24,
|
||||||
"memory_capacity": 2,
|
"memory_capacity": 2,
|
||||||
"model_load": 13,
|
"model_load": 13,
|
||||||
@@ -6131,14 +6131,14 @@
|
|||||||
"日志缺失": 1,
|
"日志缺失": 1,
|
||||||
"验证失败": 23
|
"验证失败": 23
|
||||||
},
|
},
|
||||||
"failureCount": 320,
|
"failureCount": 321,
|
||||||
"failureRate": 0.9846,
|
"failureRate": 0.9847,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 1,
|
"platformFailureCount": 1,
|
||||||
"successCount": 5,
|
"successCount": 5,
|
||||||
"successRate": 0.0154,
|
"successRate": 0.0153,
|
||||||
"total": 325,
|
"total": 326,
|
||||||
"unresolvedFailureCount": 266
|
"unresolvedFailureCount": 266
|
||||||
},
|
},
|
||||||
"Kunlunxin_r-200-8f": {
|
"Kunlunxin_r-200-8f": {
|
||||||
@@ -17678,14 +17678,14 @@
|
|||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": {
|
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": {
|
||||||
"attributableFailureCount": 2,
|
"attributableFailureCount": 3,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 2,
|
"decisionTotal": 3,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"backend_operator": 2
|
"backend_operator": 3
|
||||||
},
|
},
|
||||||
"failureCount": 2,
|
"failureCount": 3,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"modelType": "starcoder2",
|
"modelType": "starcoder2",
|
||||||
@@ -17697,7 +17697,7 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Kunlunxin_p-800",
|
"targetGpu": "Kunlunxin_p-800",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 2,
|
"total": 3,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
"MetaX_c-500|vllm|text-generation|aquila3|none": {
|
"MetaX_c-500|vllm|text-generation|aquila3|none": {
|
||||||
@@ -20888,9 +20888,9 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 16
|
"ambiguous_runtime": 15
|
||||||
},
|
},
|
||||||
"failureCount": 16,
|
"failureCount": 15,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "transformers",
|
"framework": "transformers",
|
||||||
"lastPlatformFailureAt": null,
|
"lastPlatformFailureAt": null,
|
||||||
@@ -20902,8 +20902,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Iluvatar_bi-100",
|
"targetGpu": "Iluvatar_bi-100",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 16,
|
"total": 15,
|
||||||
"unresolvedFailureCount": 16
|
"unresolvedFailureCount": 15
|
||||||
},
|
},
|
||||||
"Iluvatar_bi-150|transformers|text-generation": {
|
"Iluvatar_bi-150|transformers|text-generation": {
|
||||||
"attributableFailureCount": 5,
|
"attributableFailureCount": 5,
|
||||||
@@ -21141,18 +21141,18 @@
|
|||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800|vllm|text-generation": {
|
"Kunlunxin_p-800|vllm|text-generation": {
|
||||||
"attributableFailureCount": 6,
|
"attributableFailureCount": 7,
|
||||||
"consecutiveFailures": 6,
|
"consecutiveFailures": 7,
|
||||||
"consecutivePlatformFailures": 0,
|
"consecutivePlatformFailures": 0,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 6,
|
"decisionTotal": 7,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"backend_operator": 2,
|
"backend_operator": 3,
|
||||||
"framework_architecture_unsupported": 4,
|
"framework_architecture_unsupported": 4,
|
||||||
"参数/模板问题": 2
|
"参数/模板问题": 2
|
||||||
},
|
},
|
||||||
"failureCount": 8,
|
"failureCount": 9,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"lastPlatformFailureAt": null,
|
"lastPlatformFailureAt": null,
|
||||||
@@ -21164,7 +21164,7 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Kunlunxin_p-800",
|
"targetGpu": "Kunlunxin_p-800",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 8,
|
"total": 9,
|
||||||
"unresolvedFailureCount": 2
|
"unresolvedFailureCount": 2
|
||||||
},
|
},
|
||||||
"MetaX_c-500|vllm|text-generation": {
|
"MetaX_c-500|vllm|text-generation": {
|
||||||
@@ -23005,9 +23005,9 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 2
|
"ambiguous_runtime": 1
|
||||||
},
|
},
|
||||||
"failureCount": 2,
|
"failureCount": 1,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "transformers",
|
"framework": "transformers",
|
||||||
"lastTerminalAt": "2026-09-21T05:42:03.443562+00:00",
|
"lastTerminalAt": "2026-09-21T05:42:03.443562+00:00",
|
||||||
@@ -23020,8 +23020,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Iluvatar_bi-100",
|
"targetGpu": "Iluvatar_bi-100",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 2,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 2
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
"Iluvatar_bi-100|transformers|text-generation|mistral3|none": {
|
"Iluvatar_bi-100|transformers|text-generation|mistral3|none": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -24150,18 +24150,18 @@
|
|||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": {
|
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors": {
|
||||||
"attributableFailureCount": 2,
|
"attributableFailureCount": 3,
|
||||||
"consecutiveFailures": 2,
|
"consecutiveFailures": 3,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 2,
|
"decisionTotal": 3,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"backend_operator": 2
|
"backend_operator": 3
|
||||||
},
|
},
|
||||||
"failureCount": 2,
|
"failureCount": 3,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"lastTerminalAt": "2026-09-22T08:25:19.951625+00:00",
|
"lastTerminalAt": "2026-09-23T02:35:34.755698+00:00",
|
||||||
"modelType": "starcoder2",
|
"modelType": "starcoder2",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
@@ -24171,7 +24171,7 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Kunlunxin_p-800",
|
"targetGpu": "Kunlunxin_p-800",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 2,
|
"total": 3,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
"Mthreads_s4000|vllm|text-generation|llama|compressed-tensors": {
|
"Mthreads_s4000|vllm|text-generation|llama|compressed-tensors": {
|
||||||
@@ -41372,6 +41372,30 @@
|
|||||||
"total": 2,
|
"total": 2,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
|
"Kunlunxin_p-800|vllm|text-generation|starcoder2|compressed-tensors|34": {
|
||||||
|
"attributableFailureCount": 1,
|
||||||
|
"decisionFailureRate": 1.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 1,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"backend_operator": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"loadSizeLog2Bucket": 34,
|
||||||
|
"modelType": "starcoder2",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "compressed-tensors",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Kunlunxin_p-800",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 0
|
||||||
|
},
|
||||||
"MetaX_c-500|vllm|text-generation|aquila3|none|12": {
|
"MetaX_c-500|vllm|text-generation|aquila3|none|12": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
"decisionFailureRate": 0.0,
|
"decisionFailureRate": 0.0,
|
||||||
@@ -45802,18 +45826,18 @@
|
|||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"terminalRecords": 16795,
|
"terminalRecords": 16796,
|
||||||
"totalRecords": 17013,
|
"totalRecords": 17014,
|
||||||
"totals": {
|
"totals": {
|
||||||
"attributableFailureCount": 6012,
|
"attributableFailureCount": 6013,
|
||||||
"decisionFailureRate": 0.8622,
|
"decisionFailureRate": 0.8622,
|
||||||
"decisionSuccessRate": 0.1378,
|
"decisionSuccessRate": 0.1378,
|
||||||
"decisionTotal": 6973,
|
"decisionTotal": 6974,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 4283,
|
"ambiguous_runtime": 4283,
|
||||||
"architecture_compatibility": 212,
|
"architecture_compatibility": 212,
|
||||||
"attention_backend": 3,
|
"attention_backend": 3,
|
||||||
"backend_operator": 114,
|
"backend_operator": 115,
|
||||||
"context_length": 322,
|
"context_length": 322,
|
||||||
"framework_architecture_unsupported": 2132,
|
"framework_architecture_unsupported": 2132,
|
||||||
"memory_capacity": 1203,
|
"memory_capacity": 1203,
|
||||||
@@ -45826,19 +45850,20 @@
|
|||||||
"日志缺失": 719,
|
"日志缺失": 719,
|
||||||
"验证失败": 676
|
"验证失败": 676
|
||||||
},
|
},
|
||||||
"failureCount": 15834,
|
"failureCount": 15835,
|
||||||
"failureRate": 0.9428,
|
"failureRate": 0.9428,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 931,
|
"platformFailureCount": 931,
|
||||||
"successCount": 961,
|
"successCount": 961,
|
||||||
"successRate": 0.0572,
|
"successRate": 0.0572,
|
||||||
"total": 16795,
|
"total": 16796,
|
||||||
"unresolvedFailureCount": 8891
|
"unresolvedFailureCount": 8891
|
||||||
},
|
},
|
||||||
"warnings": [
|
"warnings": [
|
||||||
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
|
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
@@ -45849,7 +45874,6 @@
|
|||||||
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
|
||||||
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Sunrise_pt-200-x1|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Sunrise_pt-200-x1|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
@@ -45873,7 +45897,6 @@
|
|||||||
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
|
||||||
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
@@ -45887,10 +45910,11 @@
|
|||||||
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
|
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||||
]
|
]
|
||||||
},
|
},
|
||||||
"storageMode": "decision_state_only",
|
"storageMode": "decision_state_only",
|
||||||
"summarizedRecords": 17013,
|
"summarizedRecords": 17014,
|
||||||
"version": 1
|
"version": 1
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -76,6 +76,7 @@
|
|||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T04:22:26.358545+00:00", "modelId": "webAI-Official/TwIL-LM3", "modelProfile": {"architectures": ["SmolLM3ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6150235096, "estimatedRequiredGiB": 24.879, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "smollm3", "modelscopeFileSize": 22261451381, "modelscopeLicense": "other", "modelscopeParams": 3075098624, "modelscopeTags": ["license:other", "model_type:smollm3", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:smollm3", "custom_tag:twil-lm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22261451381}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:17:59.369215+00:00", "targetGpu": "Biren_166m", "taskId": "5002150", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T04:22:26.358545+00:00", "modelId": "webAI-Official/TwIL-LM3", "modelProfile": {"architectures": ["SmolLM3ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6150235096, "estimatedRequiredGiB": 24.879, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "smollm3", "modelscopeFileSize": 22261451381, "modelscopeLicense": "other", "modelscopeParams": 3075098624, "modelscopeTags": ["license:other", "model_type:smollm3", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:smollm3", "custom_tag:twil-lm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22261451381}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:17:59.369215+00:00", "targetGpu": "Biren_166m", "taskId": "5002150", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T16:20:26.367867+00:00", "modelId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-8bit", "modelProfile": {"architectures": ["MuseGlimmerForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 35620174949, "estimatedRequiredGiB": 39.84, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "muse_glimmer", "modelscopeFileSize": 35648665467, "modelscopeLicense": "apache-2.0", "modelscopeParams": 9345525760, "modelscopeTags": ["license:apache-2.0", "model_type:muse_glimmer", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:muse_glimmer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 35648665467}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T13:43:30.755806+00:00", "targetGpu": "Biren_166m", "taskId": "5001736", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T16:20:26.367867+00:00", "modelId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-8bit", "modelProfile": {"architectures": ["MuseGlimmerForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 35620174949, "estimatedRequiredGiB": 39.84, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "muse_glimmer", "modelscopeFileSize": 35648665467, "modelscopeLicense": "apache-2.0", "modelscopeParams": 9345525760, "modelscopeTags": ["license:apache-2.0", "model_type:muse_glimmer", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:muse_glimmer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 35648665467}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T13:43:30.755806+00:00", "targetGpu": "Biren_166m", "taskId": "5001736", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-23T00:07:02.456914+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:49:01.244975+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001063", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-23T00:07:02.456914+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:49:01.244975+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001063", "taskType": "text-generation", "verifyResult": null}
|
||||||
|
{"failReason": "backend_operator", "failureAction": "prefer_other_proven_framework_or_gpu", "failureCategory": "backend_operator", "failureClassificationReason": "backend_version_dependent", "failureCode": "MISSING_OPERATOR", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-23T02:35:34.755698+00:00", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663867, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663867}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:49:01.235875+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001055", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "参数/模板问题", "framework": "vllm-customized", "lastSyncTime": "2026-09-22T09:59:25.057632+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:25:50.997483+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000706", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": "参数/模板问题", "framework": "vllm-customized", "lastSyncTime": "2026-09-22T09:59:25.057632+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T12:25:50.997483+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000706", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859282+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.453427+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000553", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859282+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.453427+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000553", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859263+00:00", "modelId": "mlx-community/gemma-4-e2b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5226595850, "estimatedRequiredGiB": 5.878, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 5259113427, "modelscopeLicense": "gemma", "modelscopeParams": 1141165347, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5259113427}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.451262+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000554", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859263+00:00", "modelId": "mlx-community/gemma-4-e2b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5226595850, "estimatedRequiredGiB": 5.878, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 5259113427, "modelscopeLicense": "gemma", "modelscopeParams": 1141165347, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5259113427}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.451262+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000554", "taskType": "text-generation", "verifyResult": -1}
|
||||||
@@ -297,4 +298,3 @@
|
|||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457647+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.994993+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986861", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457647+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.994993+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986861", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_vl"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:14:26.178418+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.991772+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986860", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_vl"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:14:26.178418+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.991772+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986860", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:25:05.959914+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293474987, "estimatedRequiredGiB": 18.232, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16313701503, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16313701503}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.969433+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986859", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:25:05.959914+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293474987, "estimatedRequiredGiB": 18.232, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16313701503, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16313701503}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.969433+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986859", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T12:09:51.363808+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.967582+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986857", "taskType": "text-generation", "verifyResult": -1}
|
|
||||||
|
|||||||
File diff suppressed because it is too large
Load Diff
File diff suppressed because it is too large
Load Diff
@@ -2,24 +2,24 @@
|
|||||||
"agentVersion": "2026.09.22.1",
|
"agentVersion": "2026.09.22.1",
|
||||||
"checksums": {
|
"checksums": {
|
||||||
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
|
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
|
||||||
".modelhub_state/architecture_compatibility_blacklist.json": "cdcfe698feaf16eac366b1264402fa79c5d03a36f254171e091485de4f782fc1",
|
".modelhub_state/architecture_compatibility_blacklist.json": "d891b45df03e4a1e56761b86b5162b5d74120bec1c693327f9b2785286a71088",
|
||||||
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
||||||
".modelhub_state/market_intelligence.json": "fd79be793623b5d9f3968d56b476a9f9956ddd811ee63924f6e96f354ee72db0",
|
".modelhub_state/market_intelligence.json": "55813e4d28cd521b087a47128ba7d1842d31418b407e8f8c8a237fc016933f6b",
|
||||||
".modelhub_state/official_capabilities.json": "2ace5dc16013b6195381aad9051bdae71e5eb59ae38a8486575b22d2dbd95321",
|
".modelhub_state/official_capabilities.json": "602abbdd56c00692f3425c589e14fd39889c418bf3e08b5e63f3cb5d33439e57",
|
||||||
".modelhub_state/outcome_checkpoint.json": "207ee0fe96f1ad0125e5ee17042986e25d246e41e593ec45452569175a757808",
|
".modelhub_state/outcome_checkpoint.json": "eb01f90f8b326f3cb5ea18888009e1ae273920b6ab85f2331e475a61609a3223",
|
||||||
".modelhub_state/queue_cleanup_latest.json": "abeedea8ce654c253250bcfd506d9b5eab392b53faf9867483b77b9a5c85233b",
|
".modelhub_state/queue_cleanup_latest.json": "abeedea8ce654c253250bcfd506d9b5eab392b53faf9867483b77b9a5c85233b",
|
||||||
".modelhub_state/recent_outcomes.jsonl": "9dfc138ee2bbd5fda6efa8088e67b8fa5b977fccc750b1c9b3a524cce64d9fea",
|
".modelhub_state/recent_outcomes.jsonl": "be89f770ae4bdc0e9de6f4ee70ec569bf6dea08f70738c7d578a4ebea2bf310a",
|
||||||
".modelhub_state/recovery_active_tasks.jsonl": "9e82b697561031e34efa522cf1d8b38506deb81594f81fb548e096631229776b",
|
".modelhub_state/recovery_active_tasks.jsonl": "10159aaebad20748935c0aabde549201d4b154b6e1fc11476f67418862a9ce1c",
|
||||||
".modelhub_state/recovery_intents.jsonl": "66ef539db5857f9619c99c01c602096c0e3cc18d1fee3998238e63e06aefa88c",
|
".modelhub_state/recovery_intents.jsonl": "ddd029a6dbae152d255f2bda464b2c17d3523e15f7f6cca2ffc394735a45dc38",
|
||||||
".modelhub_state/routing_intelligence.json": "39ec7b5e77e11b96887ed828d68b094ab47d4d4e68b51e1ea92f4df7a81ab3e3",
|
".modelhub_state/routing_intelligence.json": "39ec7b5e77e11b96887ed828d68b094ab47d4d4e68b51e1ea92f4df7a81ab3e3",
|
||||||
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
||||||
".modelhub_state/worker_crashes.jsonl": "619eac16f716b2420672872ff4ce92d452f5fced4f96e8dbb37f9ce511b53c5b",
|
".modelhub_state/worker_crashes.jsonl": "619eac16f716b2420672872ff4ce92d452f5fced4f96e8dbb37f9ce511b53c5b",
|
||||||
"ledger/submissions.jsonl": "1949681660efa39796ca1a4f1338205cd3a2d7ae211571a2b78e7fc663fcba49",
|
"ledger/submissions.jsonl": "1949681660efa39796ca1a4f1338205cd3a2d7ae211571a2b78e7fc663fcba49",
|
||||||
"outcomes/submissions.jsonl": "0f7921d4c479dda8e98cb0985f6b6d1b3b1615de93ab424b1357bac650a23583"
|
"outcomes/submissions.jsonl": "860e986a50ebcb1a9c0bd01cf90dca36eeb907a901d39e16e56191a0f3f38c0e"
|
||||||
},
|
},
|
||||||
"generation": 13297,
|
"generation": 13298,
|
||||||
"phase": "cycle",
|
"phase": "cycle",
|
||||||
"schemaVersion": 1,
|
"schemaVersion": 1,
|
||||||
"updatedAt": "2026-09-23T02:34:34.092787+00:00",
|
"updatedAt": "2026-09-23T02:35:51.180630+00:00",
|
||||||
"writerId": "328f98096427428498653b568f0d5041"
|
"writerId": "328f98096427428498653b568f0d5041"
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -678,7 +678,6 @@
|
|||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063157+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020365960, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022810352, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022810352}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.259499+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001065", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063157+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020365960, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022810352, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022810352}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.259499+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001065", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063192+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020349648, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022916227, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022916227}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.241568+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001058", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:50:04.063192+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020349648, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022916227, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022916227}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.241568+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001058", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063175+00:00", "modelId": "BAAI/AquilaMed-RL", "modelProfile": {"architectures": ["AquilaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6968, "estimatedRequiredGiB": 18.386, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "aquila3", "modelscopeFileSize": 16451477782, "modelscopeLicense": "other", "modelscopeParams": 8223748096, "modelscopeTags": ["license:other", "model_type:aquila3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:arxiv:2406.12182"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16451477782}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.257127+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001059", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063175+00:00", "modelId": "BAAI/AquilaMed-RL", "modelProfile": {"architectures": ["AquilaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6968, "estimatedRequiredGiB": 18.386, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "aquila3", "modelscopeFileSize": 16451477782, "modelscopeLicense": "other", "modelscopeParams": 8223748096, "modelscopeTags": ["license:other", "model_type:aquila3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:arxiv:2406.12182"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16451477782}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.257127+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001059", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063143+00:00", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663867, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663867}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.235875+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001055", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063105+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B", "modelProfile": {"architectures": ["Lfm2MoeForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16936006912, "estimatedRequiredGiB": 18.948, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2_moe", "modelscopeFileSize": 16953948786, "modelscopeLicense": "other", "modelscopeParams": 8467856832, "modelscopeTags": ["license:other", "model_type:lfm2_moe", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16953948786}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.246640+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001066", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063105+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B", "modelProfile": {"architectures": ["Lfm2MoeForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16936006912, "estimatedRequiredGiB": 18.948, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2_moe", "modelscopeFileSize": 16953948786, "modelscopeLicense": "other", "modelscopeParams": 8467856832, "modelscopeTags": ["license:other", "model_type:lfm2_moe", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16953948786}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.246640+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001066", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063182+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.237142+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001062", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063182+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.237142+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001062", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063150+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 42709318472, "estimatedRequiredGiB": 47.753, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 42728540075, "modelscopeLicense": "llama3.1", "modelscopeParams": 70553707056, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 42728540075}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.249214+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001061", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T20:50:04.063150+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 42709318472, "estimatedRequiredGiB": 47.753, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 42728540075, "modelscopeLicense": "llama3.1", "modelscopeParams": 70553707056, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 42728540075}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T12:49:01.249214+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5001061", "taskType": "text-generation", "verifyResult": null}
|
||||||
|
|||||||
Reference in New Issue
Block a user