state: generation 19010 (cycle)
This commit is contained in:
@@ -3320,7 +3320,7 @@
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-10-01T01:02:26.749862+00:00",
|
||||
"generatedAt": "2026-10-01T01:11:14.308409+00:00",
|
||||
"summary": {
|
||||
"activeBlockCount": 167,
|
||||
"byGpuFramework": {
|
||||
|
||||
@@ -434,7 +434,7 @@
|
||||
}
|
||||
},
|
||||
"frameworkUpdatedAt": null,
|
||||
"generatedAt": "2026-10-01T01:10:09.859014+00:00",
|
||||
"generatedAt": "2026-10-01T01:13:10.264929+00:00",
|
||||
"gpuStats": {
|
||||
"Ascend_910-b3": {
|
||||
"available": true,
|
||||
|
||||
@@ -1,5 +1,5 @@
|
||||
{
|
||||
"catalogUpdatedAt": "2026-10-01T01:10:09.859014+00:00",
|
||||
"catalogUpdatedAt": "2026-10-01T01:13:10.264929+00:00",
|
||||
"configuredTaskTypes": [
|
||||
"text-generation"
|
||||
],
|
||||
@@ -56,7 +56,7 @@
|
||||
"time-series-forecasting"
|
||||
],
|
||||
"errors": [],
|
||||
"generatedAt": "2026-10-01T01:10:09.859014+00:00",
|
||||
"generatedAt": "2026-10-01T01:13:10.264929+00:00",
|
||||
"gpuCatalog": {
|
||||
"Ascend_910-b3": {
|
||||
"canVerify": true,
|
||||
@@ -6381,6 +6381,6 @@
|
||||
"updateTime": "2025-12-22 08:59:53"
|
||||
}
|
||||
],
|
||||
"taskTreeUpdatedAt": "2026-10-01T01:10:09.859014+00:00",
|
||||
"taskTreeUpdatedAt": "2026-10-01T01:13:10.264929+00:00",
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -1,6 +1,6 @@
|
||||
{
|
||||
"generatedAt": "2026-09-30T23:01:56.806499+00:00",
|
||||
"lastSyncTime": "2026-09-30T23:01:56.540688+00:00",
|
||||
"generatedAt": "2026-10-01T01:11:14.215846+00:00",
|
||||
"lastSyncTime": "2026-10-01T01:11:13.658711+00:00",
|
||||
"recentLimit": 300,
|
||||
"report": {
|
||||
"architectureCompatibilityBlocks": {
|
||||
@@ -3448,25 +3448,25 @@
|
||||
"decisionSuccessRate": 0.093,
|
||||
"decisionTotal": 43,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 77,
|
||||
"ambiguous_runtime": 79,
|
||||
"context_length": 8,
|
||||
"framework_architecture_unsupported": 30,
|
||||
"memory_capacity": 1,
|
||||
"platform_infrastructure": 1,
|
||||
"参数/模板问题": 8
|
||||
},
|
||||
"failureCount": 125,
|
||||
"failureRate": 0.969,
|
||||
"failureCount": 127,
|
||||
"failureRate": 0.9695,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 1,
|
||||
"successCount": 4,
|
||||
"successRate": 0.031,
|
||||
"successRate": 0.0305,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 129,
|
||||
"unresolvedFailureCount": 85
|
||||
"total": 131,
|
||||
"unresolvedFailureCount": 87
|
||||
},
|
||||
"Ascend_910-b3|vllm|text-generation": {
|
||||
"attributableFailureCount": 69,
|
||||
@@ -5465,18 +5465,18 @@
|
||||
"unresolvedFailureCount": 296
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": {
|
||||
"attributableFailureCount": 117,
|
||||
"attributableFailureCount": 118,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 117,
|
||||
"decisionTotal": 118,
|
||||
"failureBreakdown": {
|
||||
"backend_operator": 9,
|
||||
"framework_architecture_unsupported": 17,
|
||||
"framework_architecture_unsupported": 18,
|
||||
"platform_infrastructure": 22,
|
||||
"tokenizer_compatibility": 91,
|
||||
"参数/模板问题": 6
|
||||
},
|
||||
"failureCount": 145,
|
||||
"failureCount": 146,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"pendingCount": 0,
|
||||
@@ -5486,7 +5486,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 145,
|
||||
"total": 146,
|
||||
"unresolvedFailureCount": 6
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm|text-generation": {
|
||||
@@ -6143,28 +6143,28 @@
|
||||
"unresolvedFailureCount": 28
|
||||
},
|
||||
"vllm_fix_tokenizer": {
|
||||
"attributableFailureCount": 156,
|
||||
"decisionFailureRate": 0.9689,
|
||||
"decisionSuccessRate": 0.0311,
|
||||
"decisionTotal": 161,
|
||||
"attributableFailureCount": 157,
|
||||
"decisionFailureRate": 0.9691,
|
||||
"decisionSuccessRate": 0.0309,
|
||||
"decisionTotal": 162,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 337,
|
||||
"backend_operator": 9,
|
||||
"framework_architecture_unsupported": 37,
|
||||
"framework_architecture_unsupported": 38,
|
||||
"memory_capacity": 10,
|
||||
"model_load": 9,
|
||||
"platform_infrastructure": 22,
|
||||
"tokenizer_compatibility": 91,
|
||||
"参数/模板问题": 13
|
||||
},
|
||||
"failureCount": 528,
|
||||
"failureCount": 529,
|
||||
"failureRate": 0.9906,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 22,
|
||||
"successCount": 5,
|
||||
"successRate": 0.0094,
|
||||
"total": 533,
|
||||
"total": 534,
|
||||
"unresolvedFailureCount": 350
|
||||
},
|
||||
"vllm_tokenizer_patch": {
|
||||
@@ -6173,7 +6173,7 @@
|
||||
"decisionSuccessRate": 0.068,
|
||||
"decisionTotal": 147,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 189,
|
||||
"ambiguous_runtime": 191,
|
||||
"backend_operator": 15,
|
||||
"context_length": 8,
|
||||
"framework_architecture_unsupported": 81,
|
||||
@@ -6185,18 +6185,18 @@
|
||||
"tokenizer_compatibility": 2,
|
||||
"参数/模板问题": 29
|
||||
},
|
||||
"failureCount": 358,
|
||||
"failureRate": 0.9728,
|
||||
"failureCount": 360,
|
||||
"failureRate": 0.973,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 3,
|
||||
"successCount": 10,
|
||||
"successRate": 0.0272,
|
||||
"total": 368,
|
||||
"unresolvedFailureCount": 218
|
||||
"successRate": 0.027,
|
||||
"total": 370,
|
||||
"unresolvedFailureCount": 220
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-30T23:01:56.793295+00:00",
|
||||
"generatedAt": "2026-10-01T01:11:14.203018+00:00",
|
||||
"gpuSummaries": {
|
||||
"Ascend_910-b3": {
|
||||
"attributableFailureCount": 109,
|
||||
@@ -6204,7 +6204,7 @@
|
||||
"decisionSuccessRate": 0.1679,
|
||||
"decisionTotal": 131,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 145,
|
||||
"ambiguous_runtime": 147,
|
||||
"context_length": 8,
|
||||
"framework_architecture_unsupported": 95,
|
||||
"memory_capacity": 2,
|
||||
@@ -6215,15 +6215,15 @@
|
||||
"日志缺失": 3,
|
||||
"验证失败": 27
|
||||
},
|
||||
"failureCount": 324,
|
||||
"failureRate": 0.9364,
|
||||
"failureCount": 326,
|
||||
"failureRate": 0.9368,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 1,
|
||||
"successCount": 22,
|
||||
"successRate": 0.0636,
|
||||
"total": 346,
|
||||
"unresolvedFailureCount": 214
|
||||
"successRate": 0.0632,
|
||||
"total": 348,
|
||||
"unresolvedFailureCount": 216
|
||||
},
|
||||
"Ascend_910-b4": {
|
||||
"attributableFailureCount": 325,
|
||||
@@ -6532,16 +6532,16 @@
|
||||
"unresolvedFailureCount": 299
|
||||
},
|
||||
"Sunrise_pt-200-x1": {
|
||||
"attributableFailureCount": 753,
|
||||
"decisionFailureRate": 0.888,
|
||||
"decisionSuccessRate": 0.112,
|
||||
"decisionTotal": 848,
|
||||
"attributableFailureCount": 754,
|
||||
"decisionFailureRate": 0.8881,
|
||||
"decisionSuccessRate": 0.1119,
|
||||
"decisionTotal": 849,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 162,
|
||||
"architecture_compatibility": 44,
|
||||
"backend_operator": 37,
|
||||
"context_length": 39,
|
||||
"framework_architecture_unsupported": 162,
|
||||
"framework_architecture_unsupported": 163,
|
||||
"memory_capacity": 113,
|
||||
"platform_infrastructure": 231,
|
||||
"repository_structure": 83,
|
||||
@@ -6550,14 +6550,14 @@
|
||||
"参数/模板问题": 228,
|
||||
"验证失败": 32
|
||||
},
|
||||
"failureCount": 1406,
|
||||
"failureRate": 0.9367,
|
||||
"failureCount": 1407,
|
||||
"failureRate": 0.9368,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 231,
|
||||
"successCount": 95,
|
||||
"successRate": 0.0633,
|
||||
"total": 1501,
|
||||
"successRate": 0.0632,
|
||||
"total": 1502,
|
||||
"unresolvedFailureCount": 422
|
||||
},
|
||||
"Vastai_va16": {
|
||||
@@ -6710,9 +6710,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 11
|
||||
"ambiguous_runtime": 12
|
||||
},
|
||||
"failureCount": 11,
|
||||
"failureCount": 12,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"modelType": "gemma2",
|
||||
@@ -6724,8 +6724,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 11,
|
||||
"unresolvedFailureCount": 11
|
||||
"total": 12,
|
||||
"unresolvedFailureCount": 12
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|gemma3|compressed-tensors": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -6894,10 +6894,10 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 6,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 14,
|
||||
"ambiguous_runtime": 15,
|
||||
"context_length": 6
|
||||
},
|
||||
"failureCount": 20,
|
||||
"failureCount": 21,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"modelType": "llama",
|
||||
@@ -6909,8 +6909,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 20,
|
||||
"unresolvedFailureCount": 14
|
||||
"total": 21,
|
||||
"unresolvedFailureCount": 15
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|llama|none": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -21434,14 +21434,14 @@
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|qwen3_5_moe|none": {
|
||||
"attributableFailureCount": 2,
|
||||
"attributableFailureCount": 3,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 2,
|
||||
"decisionTotal": 3,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 2
|
||||
"framework_architecture_unsupported": 3
|
||||
},
|
||||
"failureCount": 2,
|
||||
"failureCount": 3,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"modelType": "qwen3_5_moe",
|
||||
@@ -21453,7 +21453,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 2,
|
||||
"total": 3,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|qwen3_5_mtp|none": {
|
||||
@@ -23690,11 +23690,11 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 3,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 9,
|
||||
"ambiguous_runtime": 11,
|
||||
"context_length": 2,
|
||||
"framework_architecture_unsupported": 1
|
||||
},
|
||||
"failureCount": 12,
|
||||
"failureCount": 14,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"lastPlatformFailureAt": null,
|
||||
@@ -23706,8 +23706,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 12,
|
||||
"unresolvedFailureCount": 9
|
||||
"total": 14,
|
||||
"unresolvedFailureCount": 11
|
||||
},
|
||||
"Ascend_910-b3|vllm|text-generation": {
|
||||
"attributableFailureCount": 2,
|
||||
@@ -24026,9 +24026,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 2
|
||||
"ambiguous_runtime": 1
|
||||
},
|
||||
"failureCount": 2,
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm-mlu",
|
||||
"lastPlatformFailureAt": null,
|
||||
@@ -24040,8 +24040,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Cambricon_mlu-370-x8",
|
||||
"taskType": "text-generation",
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 2
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Iluvatar_bi-150|transformers|text-generation": {
|
||||
"attributableFailureCount": 5,
|
||||
@@ -24469,11 +24469,11 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 5,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 2,
|
||||
"ambiguous_runtime": 1,
|
||||
"framework_architecture_unsupported": 4,
|
||||
"model_load": 1
|
||||
},
|
||||
"failureCount": 7,
|
||||
"failureCount": 6,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm",
|
||||
"lastPlatformFailureAt": null,
|
||||
@@ -24485,8 +24485,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "hygon_k100-ai",
|
||||
"taskType": "text-generation",
|
||||
"total": 7,
|
||||
"unresolvedFailureCount": 2
|
||||
"total": 6,
|
||||
"unresolvedFailureCount": 1
|
||||
}
|
||||
},
|
||||
"recentProfileCombinationStats": {
|
||||
@@ -24497,12 +24497,12 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 2
|
||||
"ambiguous_runtime": 3
|
||||
},
|
||||
"failureCount": 2,
|
||||
"failureCount": 3,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"lastTerminalAt": "2026-09-30T15:29:22.559055+00:00",
|
||||
"lastTerminalAt": "2026-10-01T01:11:13.658662+00:00",
|
||||
"modelType": "gemma2",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
@@ -24512,8 +24512,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 2
|
||||
"total": 3,
|
||||
"unresolvedFailureCount": 3
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|llama|compressed-tensors": {
|
||||
"attributableFailureCount": 2,
|
||||
@@ -24522,10 +24522,10 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 2,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 3,
|
||||
"ambiguous_runtime": 4,
|
||||
"context_length": 2
|
||||
},
|
||||
"failureCount": 5,
|
||||
"failureCount": 6,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"lastTerminalAt": "2026-09-30T15:29:22.559035+00:00",
|
||||
@@ -24538,8 +24538,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 5,
|
||||
"unresolvedFailureCount": 3
|
||||
"total": 6,
|
||||
"unresolvedFailureCount": 4
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|mistral|compressed-tensors": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -25691,31 +25691,6 @@
|
||||
"total": 4,
|
||||
"unresolvedFailureCount": 4
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|llama|compressed-tensors": {
|
||||
"attributableFailureCount": 0,
|
||||
"consecutiveFailures": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm-mlu",
|
||||
"lastTerminalAt": "2026-09-23T04:49:01.469123+00:00",
|
||||
"modelType": "llama",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "compressed-tensors",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Cambricon_mlu-370-x8",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|mistral|compressed-tensors": {
|
||||
"attributableFailureCount": 0,
|
||||
"consecutiveFailures": 0,
|
||||
@@ -27767,9 +27742,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 7
|
||||
"ambiguous_runtime": 8
|
||||
},
|
||||
"failureCount": 7,
|
||||
"failureCount": 8,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"loadSizeLog2Bucket": 32,
|
||||
@@ -27782,8 +27757,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 7,
|
||||
"unresolvedFailureCount": 7
|
||||
"total": 8,
|
||||
"unresolvedFailureCount": 8
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|gemma2|compressed-tensors|33": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -28224,9 +28199,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 8
|
||||
"ambiguous_runtime": 9
|
||||
},
|
||||
"failureCount": 8,
|
||||
"failureCount": 9,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"loadSizeLog2Bucket": 33,
|
||||
@@ -28239,8 +28214,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 8,
|
||||
"unresolvedFailureCount": 8
|
||||
"total": 9,
|
||||
"unresolvedFailureCount": 9
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|llama|compressed-tensors|35": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -50497,14 +50472,14 @@
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|qwen3_5_moe|none|34": {
|
||||
"attributableFailureCount": 1,
|
||||
"attributableFailureCount": 2,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"decisionTotal": 2,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 1
|
||||
"framework_architecture_unsupported": 2
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureCount": 2,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"loadSizeLog2Bucket": 34,
|
||||
@@ -50517,7 +50492,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|qwen3_5_moe|none|35": {
|
||||
@@ -53885,20 +53860,20 @@
|
||||
"unresolvedFailureCount": 0
|
||||
}
|
||||
},
|
||||
"terminalRecords": 17172,
|
||||
"totalRecords": 17393,
|
||||
"terminalRecords": 17175,
|
||||
"totalRecords": 17396,
|
||||
"totals": {
|
||||
"attributableFailureCount": 6154,
|
||||
"decisionFailureRate": 0.8632,
|
||||
"decisionSuccessRate": 0.1368,
|
||||
"decisionTotal": 7129,
|
||||
"attributableFailureCount": 6155,
|
||||
"decisionFailureRate": 0.8633,
|
||||
"decisionSuccessRate": 0.1367,
|
||||
"decisionTotal": 7130,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 4477,
|
||||
"ambiguous_runtime": 4479,
|
||||
"architecture_compatibility": 212,
|
||||
"attention_backend": 3,
|
||||
"backend_operator": 128,
|
||||
"context_length": 326,
|
||||
"framework_architecture_unsupported": 2201,
|
||||
"framework_architecture_unsupported": 2202,
|
||||
"memory_capacity": 1206,
|
||||
"model_load": 550,
|
||||
"platform_infrastructure": 943,
|
||||
@@ -53909,30 +53884,30 @@
|
||||
"日志缺失": 719,
|
||||
"验证失败": 676
|
||||
},
|
||||
"failureCount": 16197,
|
||||
"failureCount": 16200,
|
||||
"failureRate": 0.9432,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 943,
|
||||
"successCount": 975,
|
||||
"successRate": 0.0568,
|
||||
"total": 17172,
|
||||
"unresolvedFailureCount": 9100
|
||||
"total": 17175,
|
||||
"unresolvedFailureCount": 9102
|
||||
},
|
||||
"warnings": [
|
||||
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Kunlunxin_p-800 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Sunrise_pt-200-x1 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Kunlunxin_p-800|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
@@ -53977,6 +53952,6 @@
|
||||
]
|
||||
},
|
||||
"storageMode": "decision_state_only",
|
||||
"summarizedRecords": 17393,
|
||||
"summarizedRecords": 17396,
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -131,6 +131,7 @@
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-24T18:34:46.362994+00:00", "modelId": "RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w8a16", "modelProfile": {"architectures": ["MistralForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 13594088584, "estimatedRequiredGiB": 15.203, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral", "modelscopeFileSize": 13603621925, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12247782400, "modelscopeTags": ["license:apache-2.0", "model_type:mistral", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 13603621925}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:48:11.810395+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5002499", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-22T02:26:33.565856+00:00", "modelId": "RedHatAI/SmolLM-1.7B-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2014789064, "estimatedRequiredGiB": 2.255, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 2018175161, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1812039680, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2018175161}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:48:11.808469+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5002501", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-23T05:38:14.666257+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-FP8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14289052464, "estimatedRequiredGiB": 15.972, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 14291553638, "modelscopeLicense": "mit", "modelscopeParams": 13960238080, "modelscopeTags": ["license:mit", "model_type:phi3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 14291553638}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:48:11.805114+00:00", "targetGpu": "Biren_166m", "taskId": "5002503", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T01:11:13.658662+00:00", "modelId": "RedHatAI/gemma-2-2b-it-quantized.w8a8", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4385522016, "estimatedRequiredGiB": 4.926, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 4407346443, "modelscopeLicense": "gemma", "modelscopeParams": 3204165888, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4407346443}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:48:11.802242+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5002502", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T03:04:40.754950+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Instruct-MLX-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.742, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663548868, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663548868}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:48:11.800344+00:00", "targetGpu": "Biren_166m", "taskId": "5002498", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T18:49:48.654842+00:00", "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:48:11.795290+00:00", "targetGpu": "Biren_166m", "taskId": "5002497", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-22T08:01:56.366395+00:00", "modelId": "nm-testing/tinyllama-marlin24-w4a16-group128", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "c25be7d7c1bf173abe564916f8a370d9b648beb3aea2a79522523aaa6b69ee29", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 640856992, "estimatedRequiredGiB": 0.718, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 642705637, "modelscopeLicense": null, "modelscopeParams": 259844096, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 642705637}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T14:48:11.790385+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5002500", "taskType": "text-generation", "verifyResult": -1}
|
||||
@@ -191,6 +192,7 @@
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedArchitectures": ["BailingMoeV2_5ForCausalLM"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-23T17:41:18.256986+00:00", "modelId": "JANGQ-AI/Ling-2.6-flash-JANGTQ", "modelProfile": {"architectures": ["BailingMoeV2_5ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 30593642304, "estimatedRequiredGiB": 34.208, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "bailing_hybrid", "modelscopeFileSize": 30608423891, "modelscopeLicense": "mit", "modelscopeParams": null, "modelscopeTags": ["license:mit", "model_type:bailing_hybrid", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:mixture-of-experts", "custom_tag:hybrid-attention", "custom_tag:mla", "custom_tag:lightning-attention", "custom_tag:jangtq", "custom_tag:jangq-ai", "custom_tag:mlx", "custom_tag:bailing", "custom_tag:ling", "custom_tag:apple-silicon"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 30608423891}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:40:34.995530+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000087", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "backend_operator", "failureAction": "prefer_other_proven_framework_or_gpu", "failureCategory": "backend_operator", "failureClassificationReason": "backend_version_dependent", "failureCode": "MISSING_OPERATOR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-23T07:54:17.156847+00:00", "modelId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34327989512, "estimatedRequiredGiB": 38.385, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 34346039951, "modelscopeLicense": "mit", "modelscopeParams": 32762123264, "modelscopeTags": ["license:mit", "model_type:qwen3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 34346039951}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:18:21.963694+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4999765", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-30T05:13:35.747552+00:00", "modelId": "RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3SmallForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8630134648, "estimatedRequiredGiB": 9.647, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3small", "modelscopeFileSize": 8632060880, "modelscopeLicense": "mit", "modelscopeParams": 7803314176, "modelscopeTags": ["license:mit", "model_type:phi3small", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 8632060880}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:18:21.960935+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4999762", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T01:11:13.658700+00:00", "modelId": "RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084040952, "estimatedRequiredGiB": 10.163, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093261800, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "recentProfileFeedback": {"attributableFailureCount": 1, "consecutiveFailures": 1, "decisionFailureRate": 1.0, "decisionSuccessRate": 0.0, "decisionTotal": 1, "failureBreakdown": {"ambiguous_runtime": 1, "context_length": 1}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm_tokenizer_patch", "lastTerminalAt": "2026-09-21T10:33:30.477761+00:00", "modelType": "llama", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "compressed-tensors", "successCount": 0, "successRate": 0.0, "targetGpu": "Ascend_910-b3", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 1}, "repositoryOnDiskBytes": 9093261800}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:18:21.958399+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4999764", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-24T05:28:34.260385+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020365960, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022810352, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022810352}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:18:21.951854+00:00", "targetGpu": "Vastai_va16", "taskId": "4999763", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["muse_glimmer"], "framework": "vllm", "lastSyncTime": "2026-09-24T16:13:20.456571+00:00", "modelId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "modelProfile": {"architectures": ["MuseGlimmerForConditionalGeneration"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21407235494, "estimatedRequiredGiB": 23.956, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "muse_glimmer", "modelscopeFileSize": 21435726263, "modelscopeLicense": "apache-2.0", "modelscopeParams": 5792290816, "modelscopeTags": ["license:apache-2.0", "model_type:muse_glimmer", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:muse_glimmer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21435726263}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:00:43.861933+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4999589", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-23T01:54:25.558841+00:00", "modelId": "sapientinc/HRM-Text-1B", "modelProfile": {"architectures": ["HrmTextForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2365606568, "estimatedRequiredGiB": 2.651, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "hrm_text", "modelscopeFileSize": 2371660433, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1182795264, "modelscopeTags": ["license:apache-2.0", "model_type:hrm_text", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:hrm", "custom_tag:hierarchical-reasoning", "custom_tag:prefix-lm", "custom_tag:pre-alignment", "custom_tag:non-chat", "custom_tag:non-instruction-tuned"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2371660433}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:00:23.702706+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4999569", "taskType": "text-generation", "verifyResult": -1}
|
||||
@@ -296,5 +298,3 @@
|
||||
{"failReason": null, "framework": "", "lastSyncTime": "2026-09-21T04:11:57.967213+00:00", "modelId": "BAAI/CareBot_Medical_multi-llama3-8b-base", "modelProfile": {}, "outcome": "success", "status": "success", "submitTime": "2026-09-21T04:11:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4457983", "taskType": "text-generation", "verifyResult": 1}
|
||||
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-21T04:11:57.967175+00:00", "modelId": "BAAI/RoboBrain2.5-8B-NV", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T04:09:34+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969052", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-23T09:46:42.852785+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm-customized", "lastTerminalAt": "2026-09-21T02:15:17.867294+00:00", "modelType": "lfm2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 1}, "repositoryOnDiskBytes": 955963132}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T03:59:26.348415+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4993324", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-23T04:49:01.469123+00:00", "modelId": "neuralmagic/Meta-Llama-3.1-8B-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084013360, "estimatedRequiredGiB": 10.162, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093204103, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm", "custom_tag:quantized", "custom_tag:8-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093204103}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T03:31:36.734307+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4993022", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-21T03:45:50.542891+00:00", "modelId": "AI-ModelScope/granite-20b-code-base", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T03:31:21+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4080015", "taskType": "text-generation", "verifyResult": -1}
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
File diff suppressed because it is too large
Load Diff
Reference in New Issue
Block a user