state: generation 5992 (cycle)
This commit is contained in:
@@ -1007,6 +1007,25 @@
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:nemotron_h": {
|
||||
"architectureSignature": "model_type:nemotron_h",
|
||||
"architectures": [],
|
||||
"evidenceCount": 1,
|
||||
"expiresAt": "2026-10-12T02:51:21+00:00",
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"latestFailureAt": "2026-09-12T02:51:21+00:00",
|
||||
"latestSuccessfulAt": null,
|
||||
"matchType": "model_type",
|
||||
"modelType": "nemotron_h",
|
||||
"sourceModelIds": [
|
||||
"nv-community/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-NVFP4"
|
||||
],
|
||||
"sourceTaskIds": [
|
||||
"4134178"
|
||||
],
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:qwen3_5": {
|
||||
"architectureSignature": "model_type:qwen3_5",
|
||||
"architectures": [],
|
||||
@@ -1326,9 +1345,9 @@
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-12T02:55:53.615460+00:00",
|
||||
"generatedAt": "2026-09-12T02:59:04.590587+00:00",
|
||||
"summary": {
|
||||
"activeBlockCount": 67,
|
||||
"activeBlockCount": 68,
|
||||
"byGpuFramework": {
|
||||
"Ascend_910-b3|vllm": 8,
|
||||
"Ascend_910-b4|vllm": 9,
|
||||
@@ -1338,7 +1357,7 @@
|
||||
"MetaX_c-500|vllm": 7,
|
||||
"Mthreads_s4000|vllm": 7,
|
||||
"Sunrise_pt-200-x1|vllm": 1,
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 2,
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 3,
|
||||
"Vastai_va16|vllm": 12,
|
||||
"Vastai_va16|vllm_fix_tokenizer": 2,
|
||||
"hygon_k100-ai|vllm": 7,
|
||||
|
||||
@@ -35,7 +35,7 @@
|
||||
"2": {
|
||||
"complete": false,
|
||||
"lastError": "ModelHubAPIError: 系统错误",
|
||||
"listingErrors": 430,
|
||||
"listingErrors": 431,
|
||||
"nextPage": 1,
|
||||
"recordsScanned": 0,
|
||||
"uniqueRecords": 0
|
||||
@@ -102,12 +102,12 @@
|
||||
"cutoffAt": "2026-09-04T03:55:51.365685+00:00",
|
||||
"failureLogsInspected": 0,
|
||||
"mode": "incremental_decision_only",
|
||||
"nextAccountIndex": 2,
|
||||
"nextAccountIndex": 3,
|
||||
"recordsScanned": 0,
|
||||
"seenTaskIds": [],
|
||||
"startedAt": "2026-09-04T03:55:51.365685+00:00",
|
||||
"terminalRecords": 0,
|
||||
"uniqueRecords": 0,
|
||||
"updatedAt": "2026-09-12T02:55:53.587554+00:00",
|
||||
"updatedAt": "2026-09-12T02:59:04.560993+00:00",
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -416,7 +416,7 @@
|
||||
}
|
||||
},
|
||||
"frameworkUpdatedAt": null,
|
||||
"generatedAt": "2026-09-12T02:54:17.744485+00:00",
|
||||
"generatedAt": "2026-09-12T02:57:05.481546+00:00",
|
||||
"gpuStats": {
|
||||
"Ascend_910-b3": {
|
||||
"available": true,
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
@@ -1,6 +1,6 @@
|
||||
{
|
||||
"generatedAt": "2026-09-12T02:48:11.405538+00:00",
|
||||
"lastSyncTime": "2026-09-12T02:48:11.158163+00:00",
|
||||
"generatedAt": "2026-09-12T02:56:54.666019+00:00",
|
||||
"lastSyncTime": "2026-09-12T02:56:54.384007+00:00",
|
||||
"recentLimit": 300,
|
||||
"report": {
|
||||
"architectureCompatibilityBlocks": {
|
||||
@@ -1011,6 +1011,25 @@
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:nemotron_h": {
|
||||
"architectureSignature": "model_type:nemotron_h",
|
||||
"architectures": [],
|
||||
"evidenceCount": 1,
|
||||
"expiresAt": "2026-10-12T02:51:21+00:00",
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"latestFailureAt": "2026-09-12T02:51:21+00:00",
|
||||
"latestSuccessfulAt": null,
|
||||
"matchType": "model_type",
|
||||
"modelType": "nemotron_h",
|
||||
"sourceModelIds": [
|
||||
"nv-community/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-NVFP4"
|
||||
],
|
||||
"sourceTaskIds": [
|
||||
"4134178"
|
||||
],
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:qwen3_5": {
|
||||
"architectureSignature": "model_type:qwen3_5",
|
||||
"architectures": [],
|
||||
@@ -1331,7 +1350,7 @@
|
||||
}
|
||||
},
|
||||
"architectureCompatibilitySummary": {
|
||||
"activeBlockCount": 67,
|
||||
"activeBlockCount": 68,
|
||||
"byGpuFramework": {
|
||||
"Ascend_910-b3|vllm": 8,
|
||||
"Ascend_910-b4|vllm": 9,
|
||||
@@ -1341,7 +1360,7 @@
|
||||
"MetaX_c-500|vllm": 7,
|
||||
"Mthreads_s4000|vllm": 7,
|
||||
"Sunrise_pt-200-x1|vllm": 1,
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 2,
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 3,
|
||||
"Vastai_va16|vllm": 12,
|
||||
"Vastai_va16|vllm_fix_tokenizer": 2,
|
||||
"hygon_k100-ai|vllm": 7,
|
||||
@@ -1482,18 +1501,18 @@
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Ascend_910-b4|vllm|text-generation": {
|
||||
"attributableFailureCount": 37,
|
||||
"attributableFailureCount": 38,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 37,
|
||||
"decisionTotal": 38,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 33,
|
||||
"framework_architecture_unsupported": 33,
|
||||
"framework_architecture_unsupported": 34,
|
||||
"memory_capacity": 1,
|
||||
"repository_structure": 1,
|
||||
"tokenizer_compatibility": 2
|
||||
},
|
||||
"failureCount": 70,
|
||||
"failureCount": 71,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm",
|
||||
"pendingCount": 0,
|
||||
@@ -1503,7 +1522,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b4",
|
||||
"taskType": "text-generation",
|
||||
"total": 70,
|
||||
"total": 71,
|
||||
"unresolvedFailureCount": 33
|
||||
},
|
||||
"Biren_166m|unknown|text-generation": {
|
||||
@@ -2082,17 +2101,17 @@
|
||||
"unresolvedFailureCount": 33
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": {
|
||||
"attributableFailureCount": 41,
|
||||
"attributableFailureCount": 42,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 41,
|
||||
"decisionTotal": 42,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 3,
|
||||
"framework_architecture_unsupported": 4,
|
||||
"platform_infrastructure": 2,
|
||||
"tokenizer_compatibility": 38,
|
||||
"参数/模板问题": 6
|
||||
},
|
||||
"failureCount": 49,
|
||||
"failureCount": 50,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"pendingCount": 0,
|
||||
@@ -2102,22 +2121,22 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 49,
|
||||
"total": 50,
|
||||
"unresolvedFailureCount": 6
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm|text-generation": {
|
||||
"attributableFailureCount": 2,
|
||||
"attributableFailureCount": 3,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 2,
|
||||
"decisionTotal": 3,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1,
|
||||
"framework_architecture_unsupported": 1,
|
||||
"platform_infrastructure": 1,
|
||||
"tokenizer_compatibility": 1,
|
||||
"tokenizer_compatibility": 2,
|
||||
"参数/模板问题": 2
|
||||
},
|
||||
"failureCount": 6,
|
||||
"failureCount": 7,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm",
|
||||
"pendingCount": 0,
|
||||
@@ -2127,7 +2146,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 6,
|
||||
"total": 7,
|
||||
"unresolvedFailureCount": 3
|
||||
},
|
||||
"Vastai_va16|unknown|text-generation": {
|
||||
@@ -2356,30 +2375,30 @@
|
||||
"unresolvedFailureCount": 862
|
||||
},
|
||||
"vllm": {
|
||||
"attributableFailureCount": 246,
|
||||
"attributableFailureCount": 248,
|
||||
"decisionFailureRate": 0.988,
|
||||
"decisionSuccessRate": 0.012,
|
||||
"decisionTotal": 249,
|
||||
"decisionTotal": 251,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 195,
|
||||
"backend_operator": 33,
|
||||
"framework_architecture_unsupported": 172,
|
||||
"framework_architecture_unsupported": 173,
|
||||
"memory_capacity": 6,
|
||||
"model_load": 20,
|
||||
"platform_infrastructure": 1,
|
||||
"repository_structure": 5,
|
||||
"runtime_memory": 6,
|
||||
"tokenizer_compatibility": 4,
|
||||
"tokenizer_compatibility": 5,
|
||||
"参数/模板问题": 24
|
||||
},
|
||||
"failureCount": 466,
|
||||
"failureCount": 468,
|
||||
"failureRate": 0.9936,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 1,
|
||||
"successCount": 3,
|
||||
"successRate": 0.0064,
|
||||
"total": 469,
|
||||
"total": 471,
|
||||
"unresolvedFailureCount": 219
|
||||
},
|
||||
"vllm-mlu": {
|
||||
@@ -2442,26 +2461,26 @@
|
||||
"unresolvedFailureCount": 4
|
||||
},
|
||||
"vllm_fix_tokenizer": {
|
||||
"attributableFailureCount": 44,
|
||||
"attributableFailureCount": 45,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 44,
|
||||
"decisionTotal": 45,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 36,
|
||||
"framework_architecture_unsupported": 5,
|
||||
"framework_architecture_unsupported": 6,
|
||||
"model_load": 1,
|
||||
"platform_infrastructure": 2,
|
||||
"tokenizer_compatibility": 38,
|
||||
"参数/模板问题": 7
|
||||
},
|
||||
"failureCount": 89,
|
||||
"failureCount": 90,
|
||||
"failureRate": 1.0,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 2,
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"total": 89,
|
||||
"total": 90,
|
||||
"unresolvedFailureCount": 43
|
||||
},
|
||||
"vllm_tokenizer_patch": {
|
||||
@@ -2483,7 +2502,7 @@
|
||||
"unresolvedFailureCount": 1
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-12T02:48:11.401315+00:00",
|
||||
"generatedAt": "2026-09-12T02:56:54.660525+00:00",
|
||||
"gpuSummaries": {
|
||||
"Ascend_910-b3": {
|
||||
"attributableFailureCount": 26,
|
||||
@@ -2509,13 +2528,13 @@
|
||||
"unresolvedFailureCount": 63
|
||||
},
|
||||
"Ascend_910-b4": {
|
||||
"attributableFailureCount": 38,
|
||||
"decisionFailureRate": 0.95,
|
||||
"decisionSuccessRate": 0.05,
|
||||
"decisionTotal": 40,
|
||||
"attributableFailureCount": 39,
|
||||
"decisionFailureRate": 0.9512,
|
||||
"decisionSuccessRate": 0.0488,
|
||||
"decisionTotal": 41,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 34,
|
||||
"framework_architecture_unsupported": 33,
|
||||
"framework_architecture_unsupported": 34,
|
||||
"memory_capacity": 1,
|
||||
"model_load": 1,
|
||||
"repository_structure": 1,
|
||||
@@ -2523,14 +2542,14 @@
|
||||
"参数/模板问题": 12,
|
||||
"验证失败": 175
|
||||
},
|
||||
"failureCount": 259,
|
||||
"failureRate": 0.9923,
|
||||
"failureCount": 260,
|
||||
"failureRate": 0.9924,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 2,
|
||||
"successRate": 0.0077,
|
||||
"total": 261,
|
||||
"successRate": 0.0076,
|
||||
"total": 262,
|
||||
"unresolvedFailureCount": 221
|
||||
},
|
||||
"Biren_166m": {
|
||||
@@ -2741,26 +2760,26 @@
|
||||
"unresolvedFailureCount": 77
|
||||
},
|
||||
"Sunrise_pt-200-x1": {
|
||||
"attributableFailureCount": 43,
|
||||
"decisionFailureRate": 0.9773,
|
||||
"decisionSuccessRate": 0.0227,
|
||||
"decisionTotal": 44,
|
||||
"attributableFailureCount": 45,
|
||||
"decisionFailureRate": 0.9783,
|
||||
"decisionSuccessRate": 0.0217,
|
||||
"decisionTotal": 46,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1,
|
||||
"framework_architecture_unsupported": 4,
|
||||
"framework_architecture_unsupported": 5,
|
||||
"platform_infrastructure": 3,
|
||||
"tokenizer_compatibility": 39,
|
||||
"tokenizer_compatibility": 40,
|
||||
"参数/模板问题": 9,
|
||||
"验证失败": 32
|
||||
},
|
||||
"failureCount": 88,
|
||||
"failureRate": 0.9888,
|
||||
"failureCount": 90,
|
||||
"failureRate": 0.989,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 3,
|
||||
"successCount": 1,
|
||||
"successRate": 0.0112,
|
||||
"total": 89,
|
||||
"successRate": 0.011,
|
||||
"total": 91,
|
||||
"unresolvedFailureCount": 42
|
||||
},
|
||||
"Vastai_va16": {
|
||||
@@ -3976,6 +3995,29 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm|text-generation|llama|compressed-tensors": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"tokenizer_compatibility": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm",
|
||||
"modelType": "llama",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "compressed-tensors",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Vastai_va16|vllm_fix_tokenizer|text-generation|hrm_text|none": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
@@ -4204,7 +4246,7 @@
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm",
|
||||
"lastPlatformFailureAt": null,
|
||||
"lastTerminalAt": "2026-09-10T15:48:17.745912+00:00",
|
||||
"lastTerminalAt": "2026-09-12T02:56:54.384007+00:00",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
@@ -4551,22 +4593,22 @@
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": {
|
||||
"attributableFailureCount": 18,
|
||||
"consecutiveFailures": 18,
|
||||
"attributableFailureCount": 17,
|
||||
"consecutiveFailures": 17,
|
||||
"consecutivePlatformFailures": 0,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 18,
|
||||
"decisionTotal": 17,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 2,
|
||||
"platform_infrastructure": 2,
|
||||
"tokenizer_compatibility": 16
|
||||
"tokenizer_compatibility": 15
|
||||
},
|
||||
"failureCount": 20,
|
||||
"failureCount": 19,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"lastPlatformFailureAt": "2026-09-10T11:01:21+00:00",
|
||||
"lastTerminalAt": "2026-09-12T02:48:11.158114+00:00",
|
||||
"lastTerminalAt": "2026-09-12T02:56:54.383989+00:00",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 2,
|
||||
@@ -4574,7 +4616,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 20,
|
||||
"total": 19,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm|text-generation": {
|
||||
@@ -6928,6 +6970,30 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm|text-generation|llama|compressed-tensors|30": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"tokenizer_compatibility": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm",
|
||||
"loadSizeLog2Bucket": 30,
|
||||
"modelType": "llama",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "compressed-tensors",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Vastai_va16|vllm_fix_tokenizer|text-generation|hrm_text|none|31": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
@@ -7121,41 +7187,40 @@
|
||||
"unresolvedFailureCount": 0
|
||||
}
|
||||
},
|
||||
"terminalRecords": 1677,
|
||||
"totalRecords": 1764,
|
||||
"terminalRecords": 1680,
|
||||
"totalRecords": 1767,
|
||||
"totals": {
|
||||
"attributableFailureCount": 449,
|
||||
"decisionFailureRate": 0.8891,
|
||||
"decisionSuccessRate": 0.1109,
|
||||
"decisionTotal": 505,
|
||||
"attributableFailureCount": 452,
|
||||
"decisionFailureRate": 0.8898,
|
||||
"decisionSuccessRate": 0.1102,
|
||||
"decisionTotal": 508,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 394,
|
||||
"backend_operator": 40,
|
||||
"framework_architecture_unsupported": 278,
|
||||
"framework_architecture_unsupported": 280,
|
||||
"memory_capacity": 11,
|
||||
"model_load": 64,
|
||||
"platform_infrastructure": 4,
|
||||
"repository_structure": 5,
|
||||
"runtime_memory": 6,
|
||||
"tokenizer_compatibility": 45,
|
||||
"tokenizer_compatibility": 46,
|
||||
"参数/模板问题": 101,
|
||||
"验证失败": 673
|
||||
},
|
||||
"failureCount": 1621,
|
||||
"failureRate": 0.9666,
|
||||
"failureCount": 1624,
|
||||
"failureRate": 0.9667,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 4,
|
||||
"successCount": 56,
|
||||
"successRate": 0.0334,
|
||||
"total": 1677,
|
||||
"successRate": 0.0333,
|
||||
"total": 1680,
|
||||
"unresolvedFailureCount": 1168
|
||||
},
|
||||
"warnings": [
|
||||
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Sunrise_pt-200-x1 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
@@ -7164,6 +7229,7 @@
|
||||
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Sunrise_pt-200-x1 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"组合 hygon_k100-ai|vllm-patch-tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b4|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
@@ -7179,11 +7245,12 @@
|
||||
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_bi-150|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Vastai_va16|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
]
|
||||
},
|
||||
"storageMode": "decision_state_only",
|
||||
"summarizedRecords": 1764,
|
||||
"summarizedRecords": 1767,
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -14,13 +14,13 @@
|
||||
100
|
||||
],
|
||||
"accounts": 12,
|
||||
"activeScanned": 1055,
|
||||
"activeScanned": 1054,
|
||||
"ageCleanupMode": "admission_only",
|
||||
"agePolicySkipped": {
|
||||
"cleanupDisabled": true,
|
||||
"reason": "admission_only"
|
||||
},
|
||||
"architectureBlockCount": 67,
|
||||
"architectureBlockCount": 68,
|
||||
"architectureFrameworkCatalog": {
|
||||
"ascend_910-b3|text-generation": [
|
||||
"llamacpp",
|
||||
@@ -98,10 +98,10 @@
|
||||
"architectureOnly": true,
|
||||
"architecturePolicySkipped": {
|
||||
"frameworkCatalogUnknown": 87,
|
||||
"frameworkContextUnknown": 330,
|
||||
"frameworkContextUnknown": 329,
|
||||
"modelArchitectureUnknown": 150,
|
||||
"noMatchingBlock": 720,
|
||||
"partiallyBlockedFrameworkSet": 98,
|
||||
"noMatchingBlock": 718,
|
||||
"partiallyBlockedFrameworkSet": 99,
|
||||
"runningMatchedProtected": 0,
|
||||
"submissionContextMismatch": 0,
|
||||
"submissionContextUnknown": 0
|
||||
|
||||
@@ -1,3 +1,5 @@
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-12T02:56:54.384007+00:00", "modelId": "amd/Qwen2.5-7B-Instruct-onnx-ryzenai-hybrid", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:53:23+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4079108", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["nemotron_h"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:56:54.383989+00:00", "modelId": "nv-community/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-NVFP4", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:51:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4134178", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:48:11.158114+00:00", "modelId": "RedHatAI/SmolLM-360M-Instruct-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:43:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4458009", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:39:29.219599+00:00", "modelId": "RedHatAI/Qwen2.5-3B-quantized.w4a16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:39:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4100876", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["mellum"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:39:29.219629+00:00", "modelId": "JetBrains/Mellum2-12B-A2.5B-Thinking", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:35:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4152611", "taskType": "text-generation", "verifyResult": -1}
|
||||
@@ -296,5 +298,3 @@
|
||||
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T13:16:14.206275+00:00", "modelId": "neuralmagic/gemma-2-9b-it-quantized.w4a16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:57:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4107036", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:42:19.103212+00:00", "modelId": "neuralmagic/Qwen2-0.5B-Instruct-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:21:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4133653", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:42:19.103221+00:00", "modelId": "RedHatAI/Llama-2-7b-chat-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:19:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4107014", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:42:19.103172+00:00", "modelId": "iic/UEmbed-4B", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:15:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4152744", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:11:42.795229+00:00", "modelId": "RedHatAI/starcoder2-3b-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:11:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4101032", "taskType": "text-generation", "verifyResult": -1}
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
File diff suppressed because it is too large
Load Diff
@@ -1,24 +1,24 @@
|
||||
{
|
||||
"agentVersion": "2026.09.04.4",
|
||||
"checksums": {
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "011eb3c1b849bf9708f1bac75ca4324fba7ccea5b2d31e7958672b438ac07a2e",
|
||||
".modelhub_state/architecture_history_backfill.json": "c40825e111c98ca7ec5253810d8b6316b151ae07f452172b53490b317ef8b8f9",
|
||||
".modelhub_state/market_intelligence.json": "425fd2d055e24890c3dbdcf7d3c53fa3b5d026de8bcd0ea2e575042233b6db62",
|
||||
".modelhub_state/official_capabilities.json": "8369b20d67b9198cfbc98c63f085f155c2ad542fce1d7c0bb89011603d855dc4",
|
||||
".modelhub_state/outcome_checkpoint.json": "267592fed01a3ba5fb68028d166d6fe5c608ecec7c2b32b1985bc4f6099f8d69",
|
||||
".modelhub_state/queue_cleanup_latest.json": "384aaa478516df0fbe81d4231e5cdb026ef79a424962096f691840bfd3cb6931",
|
||||
".modelhub_state/recent_outcomes.jsonl": "f40a93b9569ca8a0358c678853f3510722522feb6de5c702465195362b816654",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "809d9b049c3245950f42f45c35800399af62b9854ae42055db6811592158e5b2",
|
||||
".modelhub_state/recovery_intents.jsonl": "74ebea4b0eafef707cdb83c55ad4ad7b0ac81b3545e148e9bdbf592a944c91cd",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "721e1bad1444c6c604a303eb5d9f50044d3c507e81060f0058923b4e695c5038",
|
||||
".modelhub_state/architecture_history_backfill.json": "3f2e1bd8a9f7ed99a909f2d1254ac7953d404478b10472da6b49167fd2932efc",
|
||||
".modelhub_state/market_intelligence.json": "4f0c21f848715ec0eaac426844286809622bcc6e62d0249a2139bf7a621b61e2",
|
||||
".modelhub_state/official_capabilities.json": "158704e2ce643a87d7843a5369834db9339a6c2db1c7e976c5b86eb55039ce75",
|
||||
".modelhub_state/outcome_checkpoint.json": "de8419898d771b485c1aeaedd1aa952b409d96ddd876a51daf2d8e645433c46f",
|
||||
".modelhub_state/queue_cleanup_latest.json": "dbff8ecbdcba815a2302661742ce4689464c0250a7ef3678d448f0d7a5dc07da",
|
||||
".modelhub_state/recent_outcomes.jsonl": "fea2c8180e648ab702cfb751b88efcae2b55820fdeb03576bab6e5108dbdc0a3",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "dd85494be676f6117957f3ec4e84bf7b558df9553e5945549601a22113cef06b",
|
||||
".modelhub_state/recovery_intents.jsonl": "cebd8ca56df5eb250c64975557e8591f5e20a5e8483360ade4410f572a5434e9",
|
||||
".modelhub_state/routing_intelligence.json": "316b7d02e32cf2dd3f7c0ca452f58d571a371863ac4d1c1b5193456035ce3a53",
|
||||
".modelhub_state/submission_exclusions.jsonl": "14480c58228be6b76e21e98008960b815d20d16132dcab7a11e5b449c5e2d220",
|
||||
".modelhub_state/worker_crashes.jsonl": "9693a31a13cc18a3136ff5373f9569dc2fecaa274aaf91d3b1c8a67f31fe0162",
|
||||
"ledger/submissions.jsonl": "8abcf838f382e71b9140db2901b93972a33b043c4c67a56473d00af9ccd9be3c",
|
||||
"outcomes/submissions.jsonl": "c53f6ac33f79e616e13d79bc0ead6c3f607869ad8325b1aec0fae42c5c18de72"
|
||||
"outcomes/submissions.jsonl": "89e4ed1669c9c8d079363347a986e7e020d5688c4a0a5b51c8c2e756ca106cb7"
|
||||
},
|
||||
"generation": 5991,
|
||||
"generation": 5992,
|
||||
"phase": "cycle",
|
||||
"schemaVersion": 1,
|
||||
"updatedAt": "2026-09-12T02:55:53.702813+00:00",
|
||||
"updatedAt": "2026-09-12T02:59:05.238590+00:00",
|
||||
"writerId": "e5dd59a4c84a4f21926bad61448789e2"
|
||||
}
|
||||
|
||||
@@ -81,7 +81,6 @@
|
||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-05T07:14:11.020202+00:00", "modelId": "inferencerlabs/DeepSeek-V4-Flash-MTP-DSpark-MLX", "modelProfile": {"architectures": [], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12994906286, "estimatedRequiredGiB": 14.53, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "deepseek_v4_dspark", "modelscopeFileSize": 13001289422, "modelscopeLicense": null, "modelscopeParams": 4276397927, "modelscopeTags": ["model_type:deepseek_v4_dspark", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13001289422}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:12:12.056590+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4636059", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T07:14:11.020212+00:00", "modelId": "neuralmagic/SmolLM-360M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ac9457423139c55bd60f70bc9092d1d9dddd724b01ac0d5cac27be6ed145e50e", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 504052632, "estimatedRequiredGiB": 0.567, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 507443106, "modelscopeLicense": "apache-2.0", "modelscopeParams": 409007040, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 507443106}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:12:18.663255+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4636060", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T07:14:11.020096+00:00", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5705684944, "estimatedRequiredGiB": 6.387, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5714910031, "modelscopeLicense": null, "modelscopeParams": 8031506432, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 5714910031}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:12:29.786627+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4636099", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T07:21:16.625484+00:00", "modelId": "nm-testing/tinyllama-oneshot-w8a8-static-v2", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ac9457423139c55bd60f70bc9092d1d9dddd724b01ac0d5cac27be6ed145e50e", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1231270112, "estimatedRequiredGiB": 1.378, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 1233118759, "modelscopeLicense": null, "modelscopeParams": 1100048384, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 1233118759}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:20:09.476751+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4636203", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-05T07:32:32.823940+00:00", "modelId": "XHToken/Spark-X2.5-1.7B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3415340496, "estimatedRequiredGiB": 3.834, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://cambricon.com/index.php?a=lists&c=index&catid=406&m=content"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 3430847031, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1707657216, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3430847031}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:29:59.902384+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4636365", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-05T07:32:32.823946+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-v0.4-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727938576, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://cambricon.com/index.php?a=lists&c=index&catid=406&m=content"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737300809, "modelscopeLicense": "other", "modelscopeParams": 8030261248, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737300809}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:29:59.915927+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4636369", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-05T07:32:32.823927+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-DPO-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://cambricon.com/index.php?a=lists&c=index&catid=406&m=content"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737317601, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737317601}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:29:59.925332+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4636370", "taskType": "text-generation", "verifyResult": null}
|
||||
@@ -618,8 +617,8 @@
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T01:25:57.322160+00:00", "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T17:22:24.744587+00:00", "targetGpu": "MetaX_c-500", "taskId": "4783634", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-12T02:21:54.921668+00:00", "modelId": "guaidao2/LFM2.5-2.6B-For-CTF", "modelProfile": {"architectures": [], "configFingerprint": "6647f6920340e837cb19ab6ec8bc02c4d5d8cebe4d720e359319e40aa1ebaf79", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2874779040, "estimatedRequiredGiB": 11.734, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 10499570876, "modelscopeLicense": null, "modelscopeParams": 2697198592, "modelscopeTags": ["library:gguf", "task:text-generation", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 10499570876}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:18:29.685555+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4784429", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T02:21:54.921691+00:00", "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:18:29.687478+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4784430", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T18:51:22.774816+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784981", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12999919368, "estimatedRequiredGiB": 14.555, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 13023155398, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3544073216, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:imatrix", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13023155398}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T18:51:22.770505+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784982", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T02:56:54.383907+00:00", "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:51:22.774816+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784981", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T02:56:54.383964+00:00", "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12999919368, "estimatedRequiredGiB": 14.555, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 13023155398, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3544073216, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:imatrix", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13023155398}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:51:22.770505+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784982", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-12T01:22:51.902664+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4791408", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12999919368, "estimatedRequiredGiB": 14.555, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 13023155398, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3544073216, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:imatrix", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 2}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-10T18:58:50.012615+00:00", "modelType": "qwen3_5", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 2}, "repositoryOnDiskBytes": 13023155398}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-12T01:22:51.904770+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4791409", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "nex-agi/Nex-N2.5-mini", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 70214494016, "estimatedRequiredGiB": 78.495, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 70235999301, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35107181936, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 2}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-09T14:00:25.508497+00:00", "modelType": "qwen3_5_moe", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 2}, "repositoryOnDiskBytes": 70235999301}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-12T01:22:51.908714+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4791410", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user