state: generation 13304 (intent)

This commit is contained in:
2026-09-23 02:44:28 +00:00
parent d2145b06c6
commit e59293110d
8 changed files with 145 additions and 81 deletions

View File

@@ -3044,7 +3044,7 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-23T02:39:33.342313+00:00",
"generatedAt": "2026-09-23T02:43:28.891356+00:00",
"summary": {
"activeBlockCount": 153,
"byGpuFramework": {

View File

@@ -396,7 +396,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-23T02:42:24.271501+00:00",
"generatedAt": "2026-09-23T02:43:41.289571+00:00",
"gpuStats": {
"Ascend_910-b4": {
"available": true,

View File

@@ -1,5 +1,5 @@
{
"catalogUpdatedAt": "2026-09-23T02:42:24.271501+00:00",
"catalogUpdatedAt": "2026-09-23T02:43:41.289571+00:00",
"configuredTaskTypes": [
"text-generation"
],
@@ -56,7 +56,7 @@
"time-series-forecasting"
],
"errors": [],
"generatedAt": "2026-09-23T02:42:24.271501+00:00",
"generatedAt": "2026-09-23T02:43:49.953286+00:00",
"gpuCatalog": {
"Ascend_910-b3": {
"canVerify": true,
@@ -1761,6 +1761,27 @@
],
"updatedAt": "2026-09-22T12:27:29.290285+00:00"
},
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8|2026-09-18T09:46:17+00:00|Kunlunxin_p-800": {
"taskTypes": [
"text-generation"
],
"updatedAt": "2026-09-23T02:43:49.885260+00:00"
},
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8|2026-09-18T09:46:17+00:00|MetaX_c-500": {
"taskTypes": [
"feature_emb",
"text-generation",
"visual-multi-modal"
],
"updatedAt": "2026-09-23T02:43:49.953286+00:00"
},
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8|2026-09-18T09:46:17+00:00|Mthreads_s4000": {
"taskTypes": [
"text-generation",
"visual-multi-modal"
],
"updatedAt": "2026-09-23T02:43:49.848061+00:00"
},
"https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-GGUF|2026-09-21T03:33:10+00:00|Ascend_910-b4": {
"taskTypes": [
"text-generation",
@@ -2342,19 +2363,6 @@
],
"updatedAt": "2026-09-22T16:43:42.121158+00:00"
},
"https://modelscope.cn/models/chenqi2026/7a-8elite|2026-09-22T00:58:21+00:00|Iluvatar_bi-150": {
"taskTypes": [
"asr",
"question_answering",
"reinforcement_learning",
"text-generation",
"text-to-image-generation",
"text_classification",
"vision_classification",
"visual-multi-modal"
],
"updatedAt": "2026-09-22T01:11:09.614434+00:00"
},
"https://modelscope.cn/models/chenqi2026/7a-8elite|2026-09-22T00:58:21+00:00|Vastai_va16": {
"taskTypes": [
"text-generation"
@@ -4397,28 +4405,12 @@
],
"updatedAt": "2026-09-22T03:14:36.209893+00:00"
},
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-32-v2_D|2026-09-21T18:34:41+00:00|hygon_k100-ai": {
"taskTypes": [
"text-generation",
"text-to-image-generation",
"visual-multi-modal"
],
"updatedAt": "2026-09-22T01:54:36.850717+00:00"
},
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-8-v2_D|2026-09-21T18:33:37+00:00|Vastai_va16": {
"taskTypes": [
"text-generation"
],
"updatedAt": "2026-09-22T05:08:38.677526+00:00"
},
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-8-v2_D|2026-09-21T18:33:37+00:00|hygon_k100-ai": {
"taskTypes": [
"text-generation",
"text-to-image-generation",
"visual-multi-modal"
],
"updatedAt": "2026-09-22T01:54:36.949758+00:00"
},
"https://modelscope.cn/models/nkkbr/Mini-K3-1H-mamba2-n64-g4-v1_B|2026-09-21T18:24:35+00:00|hygon_k100-ai": {
"taskTypes": [
"text-generation",
@@ -6693,6 +6685,6 @@
"updateTime": "2025-12-22 08:59:53"
}
],
"taskTreeUpdatedAt": "2026-09-23T02:42:24.271501+00:00",
"taskTreeUpdatedAt": "2026-09-23T02:43:41.289571+00:00",
"version": 1
}

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-23T02:39:33.263628+00:00",
"lastSyncTime": "2026-09-23T02:39:29.759281+00:00",
"generatedAt": "2026-09-23T02:43:28.816899+00:00",
"lastSyncTime": "2026-09-23T02:43:25.562726+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -3700,14 +3700,14 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 12,
"failureBreakdown": {
"ambiguous_runtime": 35,
"ambiguous_runtime": 36,
"framework_architecture_unsupported": 8,
"model_load": 3,
"platform_infrastructure": 2,
"tokenizer_compatibility": 1,
"参数/模板问题": 5
},
"failureCount": 54,
"failureCount": 55,
"failureRate": 1.0,
"framework": "vllm-customized",
"pendingCount": 0,
@@ -3717,8 +3717,8 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 54,
"unresolvedFailureCount": 40
"total": 55,
"unresolvedFailureCount": 41
},
"Cambricon_mlu-370-x8|vllm-mlu|text-generation": {
"attributableFailureCount": 14,
@@ -5747,22 +5747,22 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 12,
"failureBreakdown": {
"ambiguous_runtime": 35,
"ambiguous_runtime": 36,
"framework_architecture_unsupported": 8,
"model_load": 3,
"platform_infrastructure": 2,
"tokenizer_compatibility": 1,
"参数/模板问题": 5
},
"failureCount": 54,
"failureCount": 55,
"failureRate": 1.0,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 0,
"successRate": 0.0,
"total": 54,
"unresolvedFailureCount": 40
"total": 55,
"unresolvedFailureCount": 41
},
"vllm-mlu": {
"attributableFailureCount": 25,
@@ -5888,7 +5888,7 @@
"unresolvedFailureCount": 175
}
},
"generatedAt": "2026-09-23T02:39:33.251720+00:00",
"generatedAt": "2026-09-23T02:43:28.804947+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 102,
@@ -6009,7 +6009,7 @@
"decisionSuccessRate": 0.1892,
"decisionTotal": 74,
"failureBreakdown": {
"ambiguous_runtime": 142,
"ambiguous_runtime": 143,
"framework_architecture_unsupported": 48,
"memory_capacity": 4,
"model_load": 5,
@@ -6018,15 +6018,15 @@
"参数/模板问题": 41,
"验证失败": 22
},
"failureCount": 267,
"failureRate": 0.9502,
"failureCount": 268,
"failureRate": 0.9504,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 14,
"successRate": 0.0498,
"total": 281,
"unresolvedFailureCount": 205
"successRate": 0.0496,
"total": 282,
"unresolvedFailureCount": 206
},
"Iluvatar_bi-100": {
"attributableFailureCount": 194,
@@ -11321,6 +11321,29 @@
"total": 1,
"unresolvedFailureCount": 0
},
"Cambricon_mlu-370-x8|vllm-customized|text-generation|qwen3_vl|none": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-customized",
"modelType": "qwen3_vl",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x8|vllm-customized|text-generation|sophia_hybrid|none": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -20888,9 +20911,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 13
"ambiguous_runtime": 12
},
"failureCount": 13,
"failureCount": 12,
"failureRate": 1.0,
"framework": "transformers",
"lastPlatformFailureAt": null,
@@ -20902,8 +20925,8 @@
"successRate": 0.0,
"targetGpu": "Iluvatar_bi-100",
"taskType": "text-generation",
"total": 13,
"unresolvedFailureCount": 13
"total": 12,
"unresolvedFailureCount": 12
},
"Iluvatar_bi-150|transformers|text-generation": {
"attributableFailureCount": 5,
@@ -22798,6 +22821,31 @@
"total": 2,
"unresolvedFailureCount": 2
},
"Cambricon_mlu-370-x8|vllm-customized|text-generation|qwen3_vl|none": {
"attributableFailureCount": 0,
"consecutiveFailures": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-customized",
"lastTerminalAt": "2026-09-23T02:43:25.562726+00:00",
"modelType": "qwen3_vl",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x8|vllm-customized|text-generation|sophia_hybrid|none": {
"attributableFailureCount": 0,
"consecutiveFailures": 0,
@@ -22980,9 +23028,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 5
"ambiguous_runtime": 4
},
"failureCount": 5,
"failureCount": 4,
"failureRate": 1.0,
"framework": "transformers",
"lastTerminalAt": "2026-09-21T05:42:03.443548+00:00",
@@ -22995,8 +23043,8 @@
"successRate": 0.0,
"targetGpu": "Iluvatar_bi-100",
"taskType": "text-generation",
"total": 5,
"unresolvedFailureCount": 5
"total": 4,
"unresolvedFailureCount": 4
},
"Iluvatar_bi-100|transformers|text-generation|lfm2|none": {
"attributableFailureCount": 0,
@@ -32080,6 +32128,30 @@
"total": 1,
"unresolvedFailureCount": 0
},
"Cambricon_mlu-370-x8|vllm-customized|text-generation|qwen3_vl|none|34": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-customized",
"loadSizeLog2Bucket": 34,
"modelType": "qwen3_vl",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x8|vllm-customized|text-generation|sophia_hybrid|none|31": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -45825,15 +45897,15 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 16798,
"totalRecords": 17016,
"terminalRecords": 16799,
"totalRecords": 17017,
"totals": {
"attributableFailureCount": 6013,
"decisionFailureRate": 0.8622,
"decisionSuccessRate": 0.1378,
"decisionTotal": 6974,
"failureBreakdown": {
"ambiguous_runtime": 4285,
"ambiguous_runtime": 4286,
"architecture_compatibility": 212,
"attention_backend": 3,
"backend_operator": 115,
@@ -45849,15 +45921,15 @@
"日志缺失": 719,
"验证失败": 676
},
"failureCount": 15837,
"failureCount": 15838,
"failureRate": 0.9428,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 931,
"successCount": 961,
"successRate": 0.0572,
"total": 16798,
"unresolvedFailureCount": 8893
"total": 16799,
"unresolvedFailureCount": 8894
},
"warnings": [
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
@@ -45870,8 +45942,8 @@
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-100 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45896,7 +45968,6 @@
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45910,10 +45981,11 @@
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 17016,
"summarizedRecords": 17017,
"version": 1
}

View File

@@ -88,6 +88,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T20:22:53.859210+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 4}, "failureCount": 4, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-20T12:46:14.462925+00:00", "modelType": "lfm2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Iluvatar_bi-150", "taskType": "text-generation", "total": 4, "unresolvedFailureCount": 4}, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.438280+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000549", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859255+00:00", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21930054875, "estimatedRequiredGiB": 24.532, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 21950415614, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6024588416, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21950415614}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.436494+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000547", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T20:22:53.859295+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 1, "consecutiveFailures": 1, "decisionFailureRate": 1.0, "decisionSuccessRate": 0.0, "decisionTotal": 1, "failureBreakdown": {"model_load": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm_tokenizer_patch", "lastTerminalAt": "2026-09-21T04:36:00.569307+00:00", "modelType": "lfm2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Iluvatar_bi-150", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 0}, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.435532+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000550", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-23T02:43:25.562726+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.405767+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000548", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "transformers", "lastSyncTime": "2026-09-21T20:22:53.859247+00:00", "modelId": "mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9003966484, "estimatedRequiredGiB": 10.099, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 9036402901, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2463563824, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9036402901}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T12:10:32.400088+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5000555", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-22T20:46:09.364327+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:54:31.274098+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000347", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-customized", "lastSyncTime": "2026-09-22T18:31:46.771079+00:00", "modelId": "RedHatAI/starcoder2-7b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7857298896, "estimatedRequiredGiB": 8.785, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7860673720, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 7400416256, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7860673720}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T11:48:23.283753+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000251", "taskType": "text-generation", "verifyResult": -1}
@@ -297,4 +298,3 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:14:26.178403+00:00", "modelId": "Arain119/Sophia", "modelProfile": {"architectures": ["SophiaForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2226652224, "estimatedRequiredGiB": 9.768, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "sophia_hybrid", "modelscopeFileSize": 8740644030, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 1113293776, "modelscopeTags": ["license:Apache License 2.0", "model_type:sophia_hybrid", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8740644030}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:24.171826+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986869", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:42:03.443575+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293473597, "estimatedRequiredGiB": 18.233, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16314185171, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:m1", "custom_tag:m2", "custom_tag:fp16", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:chat", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b", "custom_tag:coding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16314185171}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:24.169691+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986868", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457611+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed-FP16", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683239637, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703971805, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:m1", "custom_tag:m2", "custom_tag:fp16", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703971805}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.997557+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986862", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T05:17:51.457647+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.994993+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986861", "taskType": "text-generation", "verifyResult": -1}

View File

@@ -2458,6 +2458,7 @@
{"batchId": "b9814996278e427981ad061a19fc68ec", "completedAt": "2026-09-22T19:09:59.529473+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-22T19:09:58.386064+00:00", "framework": "vllm", "intentId": "18422aa9a8834f06836946ed45997424", "lastModified": "2026-09-22T18:53:47+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-16-v2", "reason": null, "reconciledAt": "2026-09-23T02:39:48.693480+00:00", "repoId": "nkkbr/Mini-K3-1H-kda-kernel-16-v2", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "5023594", "taskType": "text-generation"}
{"batchId": "b9814996278e427981ad061a19fc68ec", "completedAt": "2026-09-22T19:09:59.529490+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-22T19:09:58.386188+00:00", "framework": "vllm", "intentId": "3d436f1c2bb84bf2a9dc7a8b998e2b65", "lastModified": "2026-09-22T18:42:06+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-2-v2", "reason": null, "reconciledAt": "2026-09-23T02:39:48.691925+00:00", "repoId": "nkkbr/Mini-K3-1H-kda-kernel-2-v2", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "5023595", "taskType": "text-generation"}
{"batchId": "b9814996278e427981ad061a19fc68ec", "completedAt": "2026-09-22T19:09:59.529495+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-22T19:09:58.386250+00:00", "framework": "vllm", "intentId": "cf52124b45d34e5eb09db6f16c17a8d7", "lastModified": "2026-09-22T18:40:12+00:00", "modelAddress": "https://modelscope.cn/models/ProCreations/BetterWright-K2-Horizon-7B-Uno", "reason": null, "reconciledAt": "2026-09-23T02:39:48.692689+00:00", "repoId": "ProCreations/BetterWright-K2-Horizon-7B-Uno", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "5023596", "taskType": "text-generation"}
{"batchId": "95bc5a65d4aa4bac9c19706aa528d821", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-23T02:44:28.518459+00:00", "framework": "vllm_fix_tokenizer", "intentId": "374addc9531642c6841513830e93c1e7", "lastModified": "2026-09-18T09:46:17+00:00", "modelAddress": "https://modelscope.cn/models/XingChen-AGI/Xing4.0-29B-A4B-FP8", "repoId": "XingChen-AGI/Xing4.0-29B-A4B-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
{"batchId": "7d0e472f09fd4e38b39dc8933c7ad231", "completedAt": "2026-09-22T13:00:50.920772+00:00", "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configSource": "modelhub_live", "createdAt": "2026-09-22T13:00:20.969224+00:00", "framework": "vllm-mlu", "intentId": "c81e5ad5eee94c7fbfbb4bd7d31b62a7", "repoId": "empero-ai/Qwen3.8-9B-Distill", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "age_policy_deferred", "targetGpu": "Cambricon_mlu-370-x4", "taskId": null, "taskType": "text-generation"}
{"batchId": "afe52d157d494615819b937d74c0906c", "completedAt": "2026-09-22T12:31:41.319873+00:00", "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configSource": "modelhub_live", "createdAt": "2026-09-22T12:31:29.032434+00:00", "framework": "vllm-mlu", "intentId": "fe2c75eab7864dca87d4316ef8729c63", "repoId": "empero-ai/Qwen3.8-9B-Distill", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "age_policy_deferred", "targetGpu": "Cambricon_mlu-370-x4", "taskId": null, "taskType": "text-generation"}
{"batchId": "2aba49d879674e828cf2c5c72b6897da", "completedAt": "2026-09-22T12:28:38.227933+00:00", "configFingerprint": "415f9524ad0e8e6c9c3d7532026f319100b3d292f6c5a09cdc37467f1144a744", "configSource": "modelhub_live", "createdAt": "2026-09-22T12:27:44.722968+00:00", "framework": "llamacpp", "intentId": "3cc699b425b64bb98e2bad28c8e2317a", "repoId": "LiquidAI/LFM2.5-2.6B-DSpark-GGUF", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Iluvatar_bi-150", "taskId": null, "taskType": "text-generation"}