state: generation 13274 (cycle)

This commit is contained in:
2026-09-23 02:03:11 +00:00
parent dda02f1ef5
commit 29bdfe0f33
9 changed files with 1695 additions and 1696 deletions

View File

@@ -3044,7 +3044,7 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-23T01:58:58.290204+00:00",
"generatedAt": "2026-09-23T02:02:54.856078+00:00",
"summary": {
"activeBlockCount": 153,
"byGpuFramework": {

View File

@@ -396,7 +396,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-23T02:01:45.345913+00:00",
"generatedAt": "2026-09-23T02:03:10.000348+00:00",
"gpuStats": {
"Ascend_910-b4": {
"available": true,

View File

@@ -1,5 +1,5 @@
{
"catalogUpdatedAt": "2026-09-23T02:01:45.345913+00:00",
"catalogUpdatedAt": "2026-09-23T02:03:10.000348+00:00",
"configuredTaskTypes": [
"text-generation"
],
@@ -56,7 +56,7 @@
"time-series-forecasting"
],
"errors": [],
"generatedAt": "2026-09-23T02:01:45.345913+00:00",
"generatedAt": "2026-09-23T02:03:10.000348+00:00",
"gpuCatalog": {
"Ascend_910-b3": {
"canVerify": true,
@@ -6669,6 +6669,6 @@
"updateTime": "2025-12-22 08:59:53"
}
],
"taskTreeUpdatedAt": "2026-09-23T02:01:45.345913+00:00",
"taskTreeUpdatedAt": "2026-09-23T02:03:10.000348+00:00",
"version": 1
}

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-23T01:58:58.220113+00:00",
"lastSyncTime": "2026-09-23T01:58:55.445357+00:00",
"generatedAt": "2026-09-23T02:02:54.781754+00:00",
"lastSyncTime": "2026-09-23T02:02:52.661615+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -4683,10 +4683,10 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 128,
"ambiguous_runtime": 129,
"参数/模板问题": 3
},
"failureCount": 131,
"failureCount": 132,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"pendingCount": 0,
@@ -4696,8 +4696,8 @@
"successRate": 0.0,
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation",
"total": 131,
"unresolvedFailureCount": 131
"total": 132,
"unresolvedFailureCount": 132
},
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation": {
"attributableFailureCount": 40,
@@ -5840,7 +5840,7 @@
"decisionSuccessRate": 0.036,
"decisionTotal": 139,
"failureBreakdown": {
"ambiguous_runtime": 297,
"ambiguous_runtime": 298,
"backend_operator": 8,
"framework_architecture_unsupported": 29,
"memory_capacity": 10,
@@ -5849,15 +5849,15 @@
"tokenizer_compatibility": 82,
"参数/模板问题": 13
},
"failureCount": 462,
"failureCount": 463,
"failureRate": 0.9893,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 18,
"successCount": 5,
"successRate": 0.0107,
"total": 467,
"unresolvedFailureCount": 310
"total": 468,
"unresolvedFailureCount": 311
},
"vllm_tokenizer_patch": {
"attributableFailureCount": 107,
@@ -5888,7 +5888,7 @@
"unresolvedFailureCount": 173
}
},
"generatedAt": "2026-09-23T01:58:58.208330+00:00",
"generatedAt": "2026-09-23T02:02:54.768381+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 102,
@@ -6119,7 +6119,7 @@
"decisionSuccessRate": 0.0862,
"decisionTotal": 58,
"failureBreakdown": {
"ambiguous_runtime": 209,
"ambiguous_runtime": 210,
"backend_operator": 10,
"framework_architecture_unsupported": 24,
"memory_capacity": 2,
@@ -6131,15 +6131,15 @@
"日志缺失": 1,
"验证失败": 23
},
"failureCount": 319,
"failureCount": 320,
"failureRate": 0.9846,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 1,
"successCount": 5,
"successRate": 0.0154,
"total": 324,
"unresolvedFailureCount": 265
"total": 325,
"unresolvedFailureCount": 266
},
"Kunlunxin_r-200-8f": {
"attributableFailureCount": 0,
@@ -16161,9 +16161,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 15
"ambiguous_runtime": 16
},
"failureCount": 15,
"failureCount": 16,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"modelType": "llama",
@@ -16175,8 +16175,8 @@
"successRate": 0.0,
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation",
"total": 15,
"unresolvedFailureCount": 15
"total": 16,
"unresolvedFailureCount": 16
},
"Kunlunxin_p-800|vllm_fix_tokenizer|text-generation|llama|none": {
"attributableFailureCount": 0,
@@ -21074,13 +21074,13 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
"ambiguous_runtime": 2
},
"failureCount": 1,
"failureCount": 2,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"lastPlatformFailureAt": null,
"lastTerminalAt": "2026-09-22T16:03:02.581784+00:00",
"lastTerminalAt": "2026-09-23T02:02:52.661615+00:00",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
@@ -21088,8 +21088,8 @@
"successRate": 0.0,
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
"total": 2,
"unresolvedFailureCount": 2
},
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation": {
"attributableFailureCount": 2,
@@ -21707,9 +21707,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 9
"ambiguous_runtime": 8
},
"failureCount": 9,
"failureCount": 8,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"lastTerminalAt": "2026-09-22T03:04:40.754950+00:00",
@@ -21722,8 +21722,8 @@
"successRate": 0.0,
"targetGpu": "Biren_166m",
"taskType": "text-generation",
"total": 9,
"unresolvedFailureCount": 9
"total": 8,
"unresolvedFailureCount": 8
},
"Biren_166m|vllm_fix_tokenizer|text-generation|llama|awq": {
"attributableFailureCount": 0,
@@ -23958,12 +23958,12 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
"ambiguous_runtime": 2
},
"failureCount": 1,
"failureCount": 2,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"lastTerminalAt": "2026-09-22T16:03:02.581784+00:00",
"lastTerminalAt": "2026-09-23T02:02:52.661615+00:00",
"modelType": "llama",
"pendingCount": 0,
"pendingRate": 0.0,
@@ -23973,8 +23973,8 @@
"successRate": 0.0,
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
"total": 2,
"unresolvedFailureCount": 2
},
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation|hrm_text|none": {
"attributableFailureCount": 1,
@@ -38979,9 +38979,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 9
"ambiguous_runtime": 10
},
"failureCount": 9,
"failureCount": 10,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"loadSizeLog2Bucket": 33,
@@ -38994,8 +38994,8 @@
"successRate": 0.0,
"targetGpu": "Kunlunxin_p-800",
"taskType": "text-generation",
"total": 9,
"unresolvedFailureCount": 9
"total": 10,
"unresolvedFailureCount": 10
},
"Kunlunxin_p-800|vllm_fix_tokenizer|text-generation|llama|none|30": {
"attributableFailureCount": 0,
@@ -45755,15 +45755,15 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 16789,
"totalRecords": 17007,
"terminalRecords": 16790,
"totalRecords": 17008,
"totals": {
"attributableFailureCount": 6011,
"decisionFailureRate": 0.8622,
"decisionSuccessRate": 0.1378,
"decisionTotal": 6972,
"failureBreakdown": {
"ambiguous_runtime": 4279,
"ambiguous_runtime": 4280,
"architecture_compatibility": 212,
"attention_backend": 3,
"backend_operator": 113,
@@ -45779,15 +45779,15 @@
"日志缺失": 719,
"验证失败": 676
},
"failureCount": 15828,
"failureCount": 15829,
"failureRate": 0.9428,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 930,
"successCount": 961,
"successRate": 0.0572,
"total": 16789,
"unresolvedFailureCount": 8887
"total": 16790,
"unresolvedFailureCount": 8888
},
"warnings": [
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
@@ -45800,8 +45800,8 @@
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-100 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45826,7 +45826,6 @@
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45840,10 +45839,11 @@
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 17007,
"summarizedRecords": 17008,
"version": 1
}

View File

@@ -160,6 +160,7 @@
{"failReason": null, "framework": "", "lastSyncTime": "2026-09-21T04:11:57.967213+00:00", "modelId": "BAAI/CareBot_Medical_multi-llama3-8b-base", "modelProfile": {}, "outcome": "success", "status": "success", "submitTime": "2026-09-21T04:11:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4457983", "taskType": "text-generation", "verifyResult": 1}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-21T04:11:57.967175+00:00", "modelId": "BAAI/RoboBrain2.5-8B-NV", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T04:09:34+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969052", "taskType": "text-generation", "verifyResult": null}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-21T03:45:50.542891+00:00", "modelId": "AI-ModelScope/granite-20b-code-base", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T03:31:21+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4080015", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-23T02:02:52.661615+00:00", "modelId": "RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084040952, "estimatedRequiredGiB": 10.163, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093261800, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093261800}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T03:11:17.186758+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4992751", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T16:03:02.581784+00:00", "modelId": "RedHatAI/SmolLM-1.7B-Instruct-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "d8df464812ff2fd3c46dc89b0fc6a96370e694fab8e52179f1ada2cdd0ea92b9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2014810120, "estimatedRequiredGiB": 2.256, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 2018195467, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1812039680, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2018195467}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T03:11:17.184850+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4992749", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "backend_operator", "failureAction": "prefer_other_proven_framework_or_gpu", "failureCategory": "backend_operator", "failureClassificationReason": "backend_version_dependent", "failureCode": "MISSING_OPERATOR", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-22T21:01:22.468534+00:00", "modelId": "neuralmagic/starcoder2-3b-FP8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3484485728, "estimatedRequiredGiB": 3.898, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 3487786187, "modelscopeLicense": "other", "modelscopeParams": 3181366272, "modelscopeTags": ["license:other", "model_type:starcoder2", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 3487786187}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T03:11:17.181043+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4992750", "taskType": "text-generation", "verifyResult": -1}
{"failReason": null, "framework": "", "lastSyncTime": "2026-09-21T03:09:23.977124+00:00", "modelId": "BAAI/Artificial-llama3_1_8B_instruct", "modelProfile": {}, "outcome": "success", "status": "success", "submitTime": "2026-09-21T03:07:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4457977", "taskType": "text-generation", "verifyResult": 1}
@@ -297,4 +298,3 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:21:25.760203+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683241105, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703487237, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703487237}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.966391+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986858", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T12:58:50.741501+00:00", "modelId": "mlx-community/gemma-4-e2b-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5222073830, "estimatedRequiredGiB": 5.872, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 5254590634, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1140034851, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5254590634}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T19:48:14.804006+00:00", "targetGpu": "Biren_166m", "taskId": "4986623", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T16:50:30.559192+00:00", "modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8674988799, "estimatedRequiredGiB": 9.718, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8695119291, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2415484144, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:speculative-decoding", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:qwen3-5", "custom_tag:9b", "custom_tag:6-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8695119291}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T19:48:14.745647+00:00", "targetGpu": "Biren_166m", "taskId": "4986622", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T18:20:35.656967+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T19:48:14.739365+00:00", "targetGpu": "Biren_166m", "taskId": "4986621", "taskType": "text-generation", "verifyResult": -1}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff