state: generation 13346 (cycle)

This commit is contained in:
2026-09-23 03:30:36 +00:00
parent 05e62652a1
commit 4f95b074a7
9 changed files with 1741 additions and 1670 deletions

View File

@@ -3044,7 +3044,7 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-23T03:26:25.550347+00:00",
"generatedAt": "2026-09-23T03:30:35.236563+00:00",
"summary": {
"activeBlockCount": 153,
"byGpuFramework": {

View File

@@ -396,7 +396,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-23T03:29:16.757350+00:00",
"generatedAt": "2026-09-23T03:30:35.254060+00:00",
"gpuStats": {
"Ascend_910-b4": {
"available": true,

View File

@@ -1,5 +1,5 @@
{
"catalogUpdatedAt": "2026-09-23T03:29:16.757350+00:00",
"catalogUpdatedAt": "2026-09-23T03:30:35.254060+00:00",
"configuredTaskTypes": [
"text-generation"
],
@@ -56,7 +56,7 @@
"time-series-forecasting"
],
"errors": [],
"generatedAt": "2026-09-23T03:29:37.952412+00:00",
"generatedAt": "2026-09-23T03:30:35.254060+00:00",
"gpuCatalog": {
"Ascend_910-b3": {
"canVerify": true,
@@ -6681,6 +6681,6 @@
"updateTime": "2025-12-22 08:59:53"
}
],
"taskTreeUpdatedAt": "2026-09-23T03:29:16.757350+00:00",
"taskTreeUpdatedAt": "2026-09-23T03:30:35.254060+00:00",
"version": 1
}

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-23T03:26:25.474628+00:00",
"lastSyncTime": "2026-09-23T03:26:25.168188+00:00",
"generatedAt": "2026-09-23T03:30:35.148806+00:00",
"lastSyncTime": "2026-09-23T03:30:34.863737+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -5065,7 +5065,7 @@
"decisionSuccessRate": 0.2371,
"decisionTotal": 97,
"failureBreakdown": {
"ambiguous_runtime": 73,
"ambiguous_runtime": 74,
"attention_backend": 2,
"backend_operator": 8,
"framework_architecture_unsupported": 53,
@@ -5077,18 +5077,18 @@
"tokenizer_compatibility": 2,
"参数/模板问题": 12
},
"failureCount": 160,
"failureRate": 0.8743,
"failureCount": 161,
"failureRate": 0.875,
"framework": "vllm",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 1,
"successCount": 23,
"successRate": 0.1257,
"successRate": 0.125,
"targetGpu": "Mthreads_s4000",
"taskType": "text-generation",
"total": 183,
"unresolvedFailureCount": 85
"total": 184,
"unresolvedFailureCount": 86
},
"Mthreads_s4000|vllm|visual-multi-modal": {
"attributableFailureCount": 1,
@@ -5717,7 +5717,7 @@
"decisionSuccessRate": 0.0259,
"decisionTotal": 3671,
"failureBreakdown": {
"ambiguous_runtime": 1534,
"ambiguous_runtime": 1535,
"architecture_compatibility": 112,
"attention_backend": 3,
"backend_operator": 87,
@@ -5731,15 +5731,15 @@
"tokenizer_compatibility": 414,
"参数/模板问题": 48
},
"failureCount": 6024,
"failureCount": 6025,
"failureRate": 0.9845,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 866,
"successCount": 95,
"successRate": 0.0155,
"total": 6119,
"unresolvedFailureCount": 1582
"total": 6120,
"unresolvedFailureCount": 1583
},
"vllm-customized": {
"attributableFailureCount": 12,
@@ -5888,7 +5888,7 @@
"unresolvedFailureCount": 175
}
},
"generatedAt": "2026-09-23T03:26:25.418062+00:00",
"generatedAt": "2026-09-23T03:30:35.136716+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 102,
@@ -6197,7 +6197,7 @@
"decisionSuccessRate": 0.3607,
"decisionTotal": 122,
"failureBreakdown": {
"ambiguous_runtime": 119,
"ambiguous_runtime": 120,
"attention_backend": 2,
"backend_operator": 8,
"framework_architecture_unsupported": 54,
@@ -6211,15 +6211,15 @@
"日志缺失": 96,
"验证失败": 26
},
"failureCount": 376,
"failureRate": 0.8952,
"failureCount": 377,
"failureRate": 0.8955,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 1,
"successCount": 44,
"successRate": 0.1048,
"total": 420,
"unresolvedFailureCount": 297
"successRate": 0.1045,
"total": 421,
"unresolvedFailureCount": 298
},
"Sunrise_pt-200-x1": {
"attributableFailureCount": 733,
@@ -18886,6 +18886,29 @@
"total": 1,
"unresolvedFailureCount": 0
},
"Mthreads_s4000|vllm|text-generation|starcoder2|compressed-tensors": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"modelType": "starcoder2",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Mthreads_s4000",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"Mthreads_s4000|vllm|text-generation|zaya|none": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -20913,9 +20936,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 12
"ambiguous_runtime": 11
},
"failureCount": 12,
"failureCount": 11,
"failureRate": 1.0,
"framework": "transformers",
"lastPlatformFailureAt": null,
@@ -20927,8 +20950,8 @@
"successRate": 0.0,
"targetGpu": "Iluvatar_bi-100",
"taskType": "text-generation",
"total": 12,
"unresolvedFailureCount": 12
"total": 11,
"unresolvedFailureCount": 11
},
"Iluvatar_bi-150|transformers|text-generation": {
"attributableFailureCount": 5,
@@ -21226,11 +21249,11 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 4,
"failureBreakdown": {
"ambiguous_runtime": 2,
"ambiguous_runtime": 3,
"framework_architecture_unsupported": 3,
"runtime_memory": 1
},
"failureCount": 6,
"failureCount": 7,
"failureRate": 1.0,
"framework": "vllm",
"lastPlatformFailureAt": null,
@@ -21242,8 +21265,8 @@
"successRate": 0.0,
"targetGpu": "Mthreads_s4000",
"taskType": "text-generation",
"total": 6,
"unresolvedFailureCount": 2
"total": 7,
"unresolvedFailureCount": 3
},
"Sunrise_pt-200-x1|unknown|text-generation": {
"attributableFailureCount": 0,
@@ -23130,9 +23153,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 2
"ambiguous_runtime": 1
},
"failureCount": 2,
"failureCount": 1,
"failureRate": 1.0,
"framework": "transformers",
"lastTerminalAt": "2026-09-21T05:42:03.443575+00:00",
@@ -23145,8 +23168,8 @@
"successRate": 0.0,
"targetGpu": "Iluvatar_bi-100",
"taskType": "text-generation",
"total": 2,
"unresolvedFailureCount": 2
"total": 1,
"unresolvedFailureCount": 1
},
"Iluvatar_bi-100|transformers|text-generation|sophia_hybrid|none": {
"attributableFailureCount": 0,
@@ -24224,6 +24247,31 @@
"total": 2,
"unresolvedFailureCount": 2
},
"Mthreads_s4000|vllm|text-generation|starcoder2|compressed-tensors": {
"attributableFailureCount": 0,
"consecutiveFailures": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"lastTerminalAt": "2026-09-23T03:30:34.863737+00:00",
"modelType": "starcoder2",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Mthreads_s4000",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|starcoder2|compressed-tensors": {
"attributableFailureCount": 0,
"consecutiveFailures": 0,
@@ -43381,6 +43429,30 @@
"total": 1,
"unresolvedFailureCount": 0
},
"Mthreads_s4000|vllm|text-generation|starcoder2|compressed-tensors|34": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"loadSizeLog2Bucket": 34,
"modelType": "starcoder2",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Mthreads_s4000",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"Mthreads_s4000|vllm|text-generation|zaya|none|32": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -45947,15 +46019,15 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 16801,
"totalRecords": 17019,
"terminalRecords": 16802,
"totalRecords": 17020,
"totals": {
"attributableFailureCount": 6015,
"decisionFailureRate": 0.8622,
"decisionSuccessRate": 0.1378,
"decisionTotal": 6976,
"failureBreakdown": {
"ambiguous_runtime": 4286,
"ambiguous_runtime": 4287,
"architecture_compatibility": 212,
"attention_backend": 3,
"backend_operator": 115,
@@ -45971,15 +46043,15 @@
"日志缺失": 719,
"验证失败": 676
},
"failureCount": 15840,
"failureCount": 15841,
"failureRate": 0.9428,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 931,
"successCount": 961,
"successRate": 0.0572,
"total": 16801,
"unresolvedFailureCount": 8894
"total": 16802,
"unresolvedFailureCount": 8895
},
"warnings": [
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
@@ -45992,8 +46064,8 @@
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -46018,6 +46090,7 @@
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -46031,11 +46104,10 @@
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 17019,
"summarizedRecords": 17020,
"version": 1
}

View File

@@ -48,6 +48,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T01:15:43.966491+00:00", "modelId": "OpenBMB/MiniCPM5-2B-MLX", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1416035216, "estimatedRequiredGiB": 1.594, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 1426233716, "modelscopeLicense": "apache-2.0", "modelscopeParams": 393390080, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1426233716}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T17:03:32.240792+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5004389", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "", "lastSyncTime": "2026-09-21T16:50:30.559368+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Quality", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T16:39:21+00:00", "targetGpu": "Biren_166m", "taskId": "4610379", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-21T16:50:30.559274+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Quality", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T16:37:21+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4332638", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-23T03:30:34.863737+00:00", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785240496, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788612744, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788612744}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T16:16:38.115066+00:00", "targetGpu": "Mthreads_s4000", "taskId": "5003582", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-23T02:39:29.759268+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020349648, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022916227, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022916227}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T16:14:31.701642+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5003541", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-22T11:22:37.063897+00:00", "modelId": "RedHatAI/Meta-Llama-3.1-8B-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084013360, "estimatedRequiredGiB": 10.162, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093203965, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm", "custom_tag:quantized", "custom_tag:8-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 3}, "failureCount": 3, "failureRate": 1.0, "framework": "vllm-customized", "lastTerminalAt": "2026-09-21T15:02:24.876374+00:00", "modelType": "llama", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "compressed-tensors", "successCount": 0, "successRate": 0.0, "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation", "total": 3, "unresolvedFailureCount": 3}, "repositoryOnDiskBytes": 9093203965}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T16:14:19.660180+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5003580", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-22T00:26:31.156547+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-5bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 804816628, "estimatedRequiredGiB": 0.905, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 809686744, "modelscopeLicense": "other", "modelscopeParams": 219544320, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 1, "consecutiveFailures": 1, "decisionFailureRate": 1.0, "decisionSuccessRate": 0.0, "decisionTotal": 1, "failureBreakdown": {"model_load": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm_tokenizer_patch", "lastTerminalAt": "2026-09-21T04:36:00.569307+00:00", "modelType": "lfm2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Iluvatar_bi-150", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 0}, "repositoryOnDiskBytes": 809686744}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T16:14:19.468977+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5003540", "taskType": "text-generation", "verifyResult": -1}
@@ -297,4 +298,3 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457633+00:00", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21930054875, "estimatedRequiredGiB": 24.532, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 21950415614, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6024588416, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21950415614}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:24.174803+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986863", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:14:26.178403+00:00", "modelId": "Arain119/Sophia", "modelProfile": {"architectures": ["SophiaForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2226652224, "estimatedRequiredGiB": 9.768, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "sophia_hybrid", "modelscopeFileSize": 8740644030, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 1113293776, "modelscopeTags": ["license:Apache License 2.0", "model_type:sophia_hybrid", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8740644030}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:24.171826+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986869", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:42:03.443575+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293473597, "estimatedRequiredGiB": 18.233, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16314185171, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:m1", "custom_tag:m2", "custom_tag:fp16", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16314185171}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:24.169691+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986868", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457611+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed-FP16", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683239637, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703971805, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:m1", "custom_tag:m2", "custom_tag:fp16", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703971805}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.997557+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986862", "taskType": "text-generation", "verifyResult": -1}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff