state: generation 13277 (cycle)

This commit is contained in:
2026-09-23 02:07:00 +00:00
parent 4bea917b75
commit dbc659adc6
9 changed files with 1692 additions and 1694 deletions

View File

@@ -3044,7 +3044,7 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-23T02:02:54.856078+00:00",
"generatedAt": "2026-09-23T02:06:47.618721+00:00",
"summary": {
"activeBlockCount": 153,
"byGpuFramework": {

View File

@@ -396,7 +396,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-23T02:05:43.076820+00:00",
"generatedAt": "2026-09-23T02:06:58.651001+00:00",
"gpuStats": {
"Ascend_910-b4": {
"available": true,

View File

@@ -1,5 +1,5 @@
{
"catalogUpdatedAt": "2026-09-23T02:05:43.076820+00:00",
"catalogUpdatedAt": "2026-09-23T02:06:58.651001+00:00",
"configuredTaskTypes": [
"text-generation"
],
@@ -56,7 +56,7 @@
"time-series-forecasting"
],
"errors": [],
"generatedAt": "2026-09-23T02:05:43.076820+00:00",
"generatedAt": "2026-09-23T02:06:58.651001+00:00",
"gpuCatalog": {
"Ascend_910-b3": {
"canVerify": true,
@@ -6669,6 +6669,6 @@
"updateTime": "2025-12-22 08:59:53"
}
],
"taskTreeUpdatedAt": "2026-09-23T02:05:43.076820+00:00",
"taskTreeUpdatedAt": "2026-09-23T02:06:58.651001+00:00",
"version": 1
}

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-23T02:02:54.781754+00:00",
"lastSyncTime": "2026-09-23T02:02:52.661615+00:00",
"generatedAt": "2026-09-23T02:06:47.498515+00:00",
"lastSyncTime": "2026-09-23T02:06:44.350388+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -3354,13 +3354,13 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 25,
"failureBreakdown": {
"ambiguous_runtime": 79,
"ambiguous_runtime": 81,
"framework_architecture_unsupported": 24,
"platform_infrastructure": 1,
"repository_structure": 1,
"参数/模板问题": 6
},
"failureCount": 111,
"failureCount": 113,
"failureRate": 1.0,
"framework": "vllm_tokenizer_patch",
"pendingCount": 0,
@@ -3370,8 +3370,8 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation",
"total": 111,
"unresolvedFailureCount": 85
"total": 113,
"unresolvedFailureCount": 87
},
"Ascend_910-b4|vllm|text-generation": {
"attributableFailureCount": 204,
@@ -5865,7 +5865,7 @@
"decisionSuccessRate": 0.0696,
"decisionTotal": 115,
"failureBreakdown": {
"ambiguous_runtime": 145,
"ambiguous_runtime": 147,
"backend_operator": 8,
"context_length": 4,
"framework_architecture_unsupported": 71,
@@ -5877,18 +5877,18 @@
"tokenizer_compatibility": 2,
"参数/模板问题": 28
},
"failureCount": 282,
"failureRate": 0.9724,
"failureCount": 284,
"failureRate": 0.9726,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 8,
"successRate": 0.0276,
"total": 290,
"unresolvedFailureCount": 173
"successRate": 0.0274,
"total": 292,
"unresolvedFailureCount": 175
}
},
"generatedAt": "2026-09-23T02:02:54.768381+00:00",
"generatedAt": "2026-09-23T02:06:47.486688+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 102,
@@ -5922,7 +5922,7 @@
"decisionSuccessRate": 0.1654,
"decisionTotal": 387,
"failureBreakdown": {
"ambiguous_runtime": 225,
"ambiguous_runtime": 227,
"context_length": 1,
"framework_architecture_unsupported": 179,
"memory_capacity": 19,
@@ -5934,15 +5934,15 @@
"日志缺失": 14,
"验证失败": 175
},
"failureCount": 1007,
"failureRate": 0.9402,
"failureCount": 1009,
"failureRate": 0.9404,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 64,
"successRate": 0.0598,
"total": 1071,
"unresolvedFailureCount": 682
"successRate": 0.0596,
"total": 1073,
"unresolvedFailureCount": 684
},
"Biren_166m": {
"attributableFailureCount": 187,
@@ -7734,9 +7734,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 22
"ambiguous_runtime": 24
},
"failureCount": 22,
"failureCount": 24,
"failureRate": 1.0,
"framework": "vllm_tokenizer_patch",
"modelType": "llama",
@@ -7748,8 +7748,8 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation",
"total": 22,
"unresolvedFailureCount": 22
"total": 24,
"unresolvedFailureCount": 24
},
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|llama|none": {
"attributableFailureCount": 0,
@@ -20635,13 +20635,13 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 8
"ambiguous_runtime": 9
},
"failureCount": 8,
"failureCount": 9,
"failureRate": 1.0,
"framework": "vllm_tokenizer_patch",
"lastPlatformFailureAt": null,
"lastTerminalAt": "2026-09-22T09:12:16.751614+00:00",
"lastTerminalAt": "2026-09-23T02:06:44.350356+00:00",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
@@ -20649,8 +20649,8 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation",
"total": 8,
"unresolvedFailureCount": 8
"total": 9,
"unresolvedFailureCount": 9
},
"Ascend_910-b4|vllm|text-generation": {
"attributableFailureCount": 0,
@@ -21406,12 +21406,12 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 7
"ambiguous_runtime": 8
},
"failureCount": 7,
"failureCount": 8,
"failureRate": 1.0,
"framework": "vllm_tokenizer_patch",
"lastTerminalAt": "2026-09-22T09:12:16.751614+00:00",
"lastTerminalAt": "2026-09-23T02:06:44.350356+00:00",
"modelType": "llama",
"pendingCount": 0,
"pendingRate": 0.0,
@@ -21421,8 +21421,8 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation",
"total": 7,
"unresolvedFailureCount": 7
"total": 8,
"unresolvedFailureCount": 8
},
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|mistral|compressed-tensors": {
"attributableFailureCount": 0,
@@ -22232,9 +22232,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 7
"ambiguous_runtime": 6
},
"failureCount": 7,
"failureCount": 6,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"lastTerminalAt": "2026-09-22T16:15:27.384896+00:00",
@@ -22247,8 +22247,8 @@
"successRate": 0.0,
"targetGpu": "Biren_166m",
"taskType": "text-generation",
"total": 7,
"unresolvedFailureCount": 7
"total": 6,
"unresolvedFailureCount": 6
},
"Biren_166m|vllm_fix_tokenizer|text-generation|qwen3|compressed-tensors": {
"attributableFailureCount": 0,
@@ -26601,9 +26601,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 11
"ambiguous_runtime": 13
},
"failureCount": 11,
"failureCount": 13,
"failureRate": 1.0,
"framework": "vllm_tokenizer_patch",
"loadSizeLog2Bucket": 33,
@@ -26616,8 +26616,8 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation",
"total": 11,
"unresolvedFailureCount": 11
"total": 13,
"unresolvedFailureCount": 13
},
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|llama|none|30": {
"attributableFailureCount": 0,
@@ -45755,15 +45755,15 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 16790,
"totalRecords": 17008,
"terminalRecords": 16792,
"totalRecords": 17010,
"totals": {
"attributableFailureCount": 6011,
"decisionFailureRate": 0.8622,
"decisionSuccessRate": 0.1378,
"decisionTotal": 6972,
"failureBreakdown": {
"ambiguous_runtime": 4280,
"ambiguous_runtime": 4282,
"architecture_compatibility": 212,
"attention_backend": 3,
"backend_operator": 113,
@@ -45779,19 +45779,20 @@
"日志缺失": 719,
"验证失败": 676
},
"failureCount": 15829,
"failureCount": 15831,
"failureRate": 0.9428,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 930,
"successCount": 961,
"successRate": 0.0572,
"total": 16790,
"unresolvedFailureCount": 8888
"total": 16792,
"unresolvedFailureCount": 8890
},
"warnings": [
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
@@ -45802,7 +45803,6 @@
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-100 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Iluvatar_bi-150|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45826,6 +45826,7 @@
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -45839,11 +45840,10 @@
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 17008,
"summarizedRecords": 17010,
"version": 1
}

View File

@@ -18,6 +18,7 @@
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-22T01:23:49.082807+00:00", "modelId": "mlx-community/Ornith-1.0-35B-OptiQ-4bit-REAP-19B", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-22T01:23:32+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4881076", "taskType": "text-generation", "verifyResult": null}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-22T01:23:49.082789+00:00", "modelId": "mlx-community/Ornith-1.5-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-22T01:23:31+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4925021", "taskType": "text-generation", "verifyResult": null}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-22T04:16:12.463726+00:00", "modelId": "RedHatAI/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663591, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663591}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T18:32:49.780633+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5005549", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-23T02:06:44.350356+00:00", "modelId": "nm-testing/Meta-llama3-8b-Instruct-SmoothQuant-Fp8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9081287016, "estimatedRequiredGiB": 10.159, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9090490505, "modelscopeLicense": null, "modelscopeParams": 8030261248, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm_tokenizer_patch", "lastTerminalAt": "2026-09-21T15:02:24.876390+00:00", "modelType": "llama", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "compressed-tensors", "successCount": 0, "successRate": 0.0, "targetGpu": "Ascend_910-b4", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 1}, "repositoryOnDiskBytes": 9090490505}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T18:22:48.749631+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5005403", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-22T02:26:33.565863+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663410366, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 1, "consecutiveFailures": 1, "decisionFailureRate": 1.0, "decisionSuccessRate": 0.0, "decisionTotal": 1, "failureBreakdown": {"model_load": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm_tokenizer_patch", "lastTerminalAt": "2026-09-21T04:36:00.569307+00:00", "modelType": "lfm2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Iluvatar_bi-150", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 0}, "repositoryOnDiskBytes": 663410366}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T18:14:31.787220+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5005326", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-22T02:26:33.565788+00:00", "modelId": "ornith-ai/Ornith-1.5-35B-A3B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "3b6ff6a21da290a9b1ca3ac432aa93c3d06a816b4b694cc9a45180f81f107aa5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23420419700, "estimatedRequiredGiB": 26.214, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 23455515266, "modelscopeLicense": "mit", "modelscopeParams": 19528501104, "modelscopeTags": ["license:mit", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 23455515266}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T18:14:20.635969+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5005325", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-22T02:26:33.565816+00:00", "modelId": "RedHatAI/starcoder2-7b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7857298896, "estimatedRequiredGiB": 8.785, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7860673720, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 7400416256, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7860673720}, "outcome": "failed", "status": "success", "submitTime": "2026-09-21T18:13:42.077548+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5005304", "taskType": "text-generation", "verifyResult": -1}
@@ -297,4 +298,3 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-21T12:09:51.363808+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.967582+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986857", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "transformers", "lastSyncTime": "2026-09-21T05:21:25.760203+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "3086a71adcdca90005b4967eb30749c44cd82389aff8d599758e4b91675d40dc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683241105, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703487237, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703487237}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T20:09:23.966391+00:00", "targetGpu": "Iluvatar_bi-100", "taskId": "4986858", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T12:58:50.741501+00:00", "modelId": "mlx-community/gemma-4-e2b-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5222073830, "estimatedRequiredGiB": 5.872, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 5254590634, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1140034851, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5254590634}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T19:48:14.804006+00:00", "targetGpu": "Biren_166m", "taskId": "4986623", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T16:50:30.559192+00:00", "modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8674988799, "estimatedRequiredGiB": 9.718, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8695119291, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2415484144, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:speculative-decoding", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:qwen3-5", "custom_tag:9b", "custom_tag:6-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8695119291}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T19:48:14.745647+00:00", "targetGpu": "Biren_166m", "taskId": "4986622", "taskType": "text-generation", "verifyResult": -1}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff