state: generation 10613 (cycle)
This commit is contained in:
@@ -1941,7 +1941,7 @@
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-20T13:59:25.770029+00:00",
|
||||
"generatedAt": "2026-09-20T14:14:40.032009+00:00",
|
||||
"summary": {
|
||||
"activeBlockCount": 94,
|
||||
"byGpuFramework": {
|
||||
|
||||
@@ -425,7 +425,7 @@
|
||||
}
|
||||
},
|
||||
"frameworkUpdatedAt": null,
|
||||
"generatedAt": "2026-09-20T14:13:36.095361+00:00",
|
||||
"generatedAt": "2026-09-20T14:14:46.591633+00:00",
|
||||
"gpuStats": {
|
||||
"Ascend_910-b3": {
|
||||
"available": true,
|
||||
|
||||
@@ -1,5 +1,5 @@
|
||||
{
|
||||
"catalogUpdatedAt": "2026-09-20T14:13:36.095361+00:00",
|
||||
"catalogUpdatedAt": "2026-09-20T14:14:46.591633+00:00",
|
||||
"configuredTaskTypes": [
|
||||
"text-generation"
|
||||
],
|
||||
@@ -56,7 +56,7 @@
|
||||
"time-series-forecasting"
|
||||
],
|
||||
"errors": [],
|
||||
"generatedAt": "2026-09-20T14:13:36.095361+00:00",
|
||||
"generatedAt": "2026-09-20T14:14:46.591633+00:00",
|
||||
"gpuCatalog": {
|
||||
"Ascend_910-b3": {
|
||||
"canVerify": true,
|
||||
@@ -6359,6 +6359,6 @@
|
||||
"updateTime": "2025-12-22 08:59:53"
|
||||
}
|
||||
],
|
||||
"taskTreeUpdatedAt": "2026-09-20T14:13:36.095361+00:00",
|
||||
"taskTreeUpdatedAt": "2026-09-20T14:14:46.591633+00:00",
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -1,6 +1,6 @@
|
||||
{
|
||||
"generatedAt": "2026-09-20T13:59:52.323200+00:00",
|
||||
"lastSyncTime": "2026-09-20T13:59:52.266077+00:00",
|
||||
"generatedAt": "2026-09-20T14:14:39.977491+00:00",
|
||||
"lastSyncTime": "2026-09-20T14:14:39.678936+00:00",
|
||||
"recentLimit": 300,
|
||||
"report": {
|
||||
"architectureCompatibilityBlocks": {
|
||||
@@ -2893,13 +2893,13 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 12,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 2,
|
||||
"ambiguous_runtime": 3,
|
||||
"framework_architecture_unsupported": 2,
|
||||
"memory_capacity": 4,
|
||||
"model_load": 2,
|
||||
"tokenizer_compatibility": 4
|
||||
},
|
||||
"failureCount": 14,
|
||||
"failureCount": 15,
|
||||
"failureRate": 1.0,
|
||||
"framework": "transformers",
|
||||
"pendingCount": 0,
|
||||
@@ -2909,8 +2909,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Iluvatar_bi-150",
|
||||
"taskType": "text-generation",
|
||||
"total": 14,
|
||||
"unresolvedFailureCount": 2
|
||||
"total": 15,
|
||||
"unresolvedFailureCount": 3
|
||||
},
|
||||
"Iluvatar_bi-150|unknown|asr": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -3415,25 +3415,25 @@
|
||||
},
|
||||
"Iluvatar_mrv-100|vllm|text-generation": {
|
||||
"attributableFailureCount": 14,
|
||||
"decisionFailureRate": 0.6087,
|
||||
"decisionSuccessRate": 0.3913,
|
||||
"decisionTotal": 23,
|
||||
"decisionFailureRate": 0.56,
|
||||
"decisionSuccessRate": 0.44,
|
||||
"decisionTotal": 25,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 4,
|
||||
"framework_architecture_unsupported": 13,
|
||||
"model_load": 1
|
||||
},
|
||||
"failureCount": 18,
|
||||
"failureRate": 0.6667,
|
||||
"failureRate": 0.6207,
|
||||
"framework": "vllm",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 9,
|
||||
"successRate": 0.3333,
|
||||
"successCount": 11,
|
||||
"successRate": 0.3793,
|
||||
"targetGpu": "Iluvatar_mrv-100",
|
||||
"taskType": "text-generation",
|
||||
"total": 27,
|
||||
"total": 29,
|
||||
"unresolvedFailureCount": 4
|
||||
},
|
||||
"Kunlunxin_p-800|unknown|text-generation": {
|
||||
@@ -4410,21 +4410,21 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 13,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 37,
|
||||
"ambiguous_runtime": 38,
|
||||
"framework_architecture_unsupported": 3,
|
||||
"memory_capacity": 4,
|
||||
"model_load": 2,
|
||||
"tokenizer_compatibility": 4
|
||||
},
|
||||
"failureCount": 50,
|
||||
"failureCount": 51,
|
||||
"failureRate": 1.0,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"total": 50,
|
||||
"unresolvedFailureCount": 37
|
||||
"total": 51,
|
||||
"unresolvedFailureCount": 38
|
||||
},
|
||||
"unknown": {
|
||||
"attributableFailureCount": 1856,
|
||||
@@ -4459,9 +4459,9 @@
|
||||
},
|
||||
"vllm": {
|
||||
"attributableFailureCount": 3474,
|
||||
"decisionFailureRate": 0.9789,
|
||||
"decisionSuccessRate": 0.0211,
|
||||
"decisionTotal": 3549,
|
||||
"decisionFailureRate": 0.9783,
|
||||
"decisionSuccessRate": 0.0217,
|
||||
"decisionTotal": 3551,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1429,
|
||||
"architecture_compatibility": 112,
|
||||
@@ -4478,13 +4478,13 @@
|
||||
"参数/模板问题": 40
|
||||
},
|
||||
"failureCount": 5807,
|
||||
"failureRate": 0.9872,
|
||||
"failureRate": 0.9869,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 864,
|
||||
"successCount": 75,
|
||||
"successRate": 0.0128,
|
||||
"total": 5882,
|
||||
"successCount": 77,
|
||||
"successRate": 0.0131,
|
||||
"total": 5884,
|
||||
"unresolvedFailureCount": 1469
|
||||
},
|
||||
"vllm-customized": {
|
||||
@@ -4618,7 +4618,7 @@
|
||||
"unresolvedFailureCount": 24
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-20T13:59:52.315429+00:00",
|
||||
"generatedAt": "2026-09-20T14:14:39.968836+00:00",
|
||||
"gpuSummaries": {
|
||||
"Ascend_910-b3": {
|
||||
"attributableFailureCount": 75,
|
||||
@@ -4785,7 +4785,7 @@
|
||||
"decisionSuccessRate": 0.1701,
|
||||
"decisionTotal": 876,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 602,
|
||||
"ambiguous_runtime": 603,
|
||||
"architecture_compatibility": 15,
|
||||
"backend_operator": 11,
|
||||
"context_length": 29,
|
||||
@@ -4800,21 +4800,21 @@
|
||||
"日志缺失": 155,
|
||||
"验证失败": 30
|
||||
},
|
||||
"failureCount": 1854,
|
||||
"failureCount": 1855,
|
||||
"failureRate": 0.9256,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 21,
|
||||
"successCount": 149,
|
||||
"successRate": 0.0744,
|
||||
"total": 2003,
|
||||
"unresolvedFailureCount": 1106
|
||||
"total": 2004,
|
||||
"unresolvedFailureCount": 1107
|
||||
},
|
||||
"Iluvatar_mrv-100": {
|
||||
"attributableFailureCount": 964,
|
||||
"decisionFailureRate": 0.8501,
|
||||
"decisionSuccessRate": 0.1499,
|
||||
"decisionTotal": 1134,
|
||||
"decisionFailureRate": 0.8486,
|
||||
"decisionSuccessRate": 0.1514,
|
||||
"decisionTotal": 1136,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 344,
|
||||
"architecture_compatibility": 47,
|
||||
@@ -4832,13 +4832,13 @@
|
||||
"验证失败": 47
|
||||
},
|
||||
"failureCount": 1684,
|
||||
"failureRate": 0.9083,
|
||||
"failureRate": 0.9073,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 7,
|
||||
"successCount": 170,
|
||||
"successRate": 0.0917,
|
||||
"total": 1854,
|
||||
"successCount": 172,
|
||||
"successRate": 0.0927,
|
||||
"total": 1856,
|
||||
"unresolvedFailureCount": 713
|
||||
},
|
||||
"Kunlunxin_p-800": {
|
||||
@@ -7094,9 +7094,10 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1,
|
||||
"model_load": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureCount": 2,
|
||||
"failureRate": 1.0,
|
||||
"framework": "transformers",
|
||||
"modelType": "mpt",
|
||||
@@ -7108,8 +7109,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Iluvatar_bi-150",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Iluvatar_bi-150|transformers|text-generation|qwen3_5|none": {
|
||||
"attributableFailureCount": 5,
|
||||
@@ -8450,23 +8451,23 @@
|
||||
"attributableFailureCount": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
"decisionSuccessRate": 1.0,
|
||||
"decisionTotal": 2,
|
||||
"decisionTotal": 4,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 0.3333,
|
||||
"failureRate": 0.2,
|
||||
"framework": "vllm",
|
||||
"modelType": "llama",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 2,
|
||||
"successRate": 0.6667,
|
||||
"successCount": 4,
|
||||
"successRate": 0.8,
|
||||
"targetGpu": "Iluvatar_mrv-100",
|
||||
"taskType": "text-generation",
|
||||
"total": 3,
|
||||
"total": 5,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Iluvatar_mrv-100|vllm|text-generation|muse_glimmer|none": {
|
||||
@@ -11240,13 +11241,13 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 12,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 2,
|
||||
"ambiguous_runtime": 3,
|
||||
"framework_architecture_unsupported": 2,
|
||||
"memory_capacity": 4,
|
||||
"model_load": 2,
|
||||
"tokenizer_compatibility": 4
|
||||
},
|
||||
"failureCount": 14,
|
||||
"failureCount": 15,
|
||||
"failureRate": 1.0,
|
||||
"framework": "transformers",
|
||||
"lastPlatformFailureAt": null,
|
||||
@@ -11258,8 +11259,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Iluvatar_bi-150",
|
||||
"taskType": "text-generation",
|
||||
"total": 14,
|
||||
"unresolvedFailureCount": 2
|
||||
"total": 15,
|
||||
"unresolvedFailureCount": 3
|
||||
},
|
||||
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation": {
|
||||
"attributableFailureCount": 9,
|
||||
@@ -11559,22 +11560,22 @@
|
||||
"decisionTotal": 9,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 3,
|
||||
"platform_infrastructure": 4,
|
||||
"platform_infrastructure": 3,
|
||||
"tokenizer_compatibility": 6
|
||||
},
|
||||
"failureCount": 13,
|
||||
"failureCount": 12,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_fix_tokenizer",
|
||||
"lastPlatformFailureAt": "2026-09-16T02:03:21+00:00",
|
||||
"lastTerminalAt": "2026-09-19T23:13:47.695526+00:00",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 4,
|
||||
"platformFailureCount": 3,
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Sunrise_pt-200-x1",
|
||||
"taskType": "text-generation",
|
||||
"total": 13,
|
||||
"total": 12,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Sunrise_pt-200-x1|vllm|text-generation": {
|
||||
@@ -11967,12 +11968,13 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1,
|
||||
"model_load": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureCount": 2,
|
||||
"failureRate": 1.0,
|
||||
"framework": "transformers",
|
||||
"lastTerminalAt": "2026-09-20T12:01:00.670989+00:00",
|
||||
"lastTerminalAt": "2026-09-20T14:14:39.678925+00:00",
|
||||
"modelType": "mpt",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
@@ -11982,8 +11984,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Iluvatar_bi-150",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Iluvatar_bi-150|transformers|text-generation|qwen3_5|none": {
|
||||
"attributableFailureCount": 5,
|
||||
@@ -15904,9 +15906,10 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1,
|
||||
"model_load": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureCount": 2,
|
||||
"failureRate": 1.0,
|
||||
"framework": "transformers",
|
||||
"loadSizeLog2Bucket": 33,
|
||||
@@ -15919,8 +15922,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Iluvatar_bi-150",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Iluvatar_bi-150|transformers|text-generation|qwen3_5|none|33": {
|
||||
"attributableFailureCount": 3,
|
||||
@@ -17908,7 +17911,7 @@
|
||||
"attributableFailureCount": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
"decisionSuccessRate": 1.0,
|
||||
"decisionTotal": 2,
|
||||
"decisionTotal": 4,
|
||||
"failureBreakdown": {},
|
||||
"failureCount": 0,
|
||||
"failureRate": 0.0,
|
||||
@@ -17919,11 +17922,11 @@
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 2,
|
||||
"successCount": 4,
|
||||
"successRate": 1.0,
|
||||
"targetGpu": "Iluvatar_mrv-100",
|
||||
"taskType": "text-generation",
|
||||
"total": 2,
|
||||
"total": 4,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Iluvatar_mrv-100|vllm|text-generation|muse_glimmer|none|34": {
|
||||
@@ -21939,15 +21942,15 @@
|
||||
"unresolvedFailureCount": 0
|
||||
}
|
||||
},
|
||||
"terminalRecords": 15859,
|
||||
"totalRecords": 15964,
|
||||
"terminalRecords": 15862,
|
||||
"totalRecords": 15967,
|
||||
"totals": {
|
||||
"attributableFailureCount": 5754,
|
||||
"decisionFailureRate": 0.8619,
|
||||
"decisionSuccessRate": 0.1381,
|
||||
"decisionTotal": 6676,
|
||||
"decisionFailureRate": 0.8616,
|
||||
"decisionSuccessRate": 0.1384,
|
||||
"decisionTotal": 6678,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 3759,
|
||||
"ambiguous_runtime": 3760,
|
||||
"architecture_compatibility": 212,
|
||||
"attention_backend": 1,
|
||||
"backend_operator": 102,
|
||||
@@ -21963,15 +21966,15 @@
|
||||
"日志缺失": 719,
|
||||
"验证失败": 673
|
||||
},
|
||||
"failureCount": 14937,
|
||||
"failureRate": 0.9419,
|
||||
"failureCount": 14938,
|
||||
"failureRate": 0.9417,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 922,
|
||||
"successCount": 922,
|
||||
"successRate": 0.0581,
|
||||
"total": 15859,
|
||||
"unresolvedFailureCount": 8261
|
||||
"successCount": 924,
|
||||
"successRate": 0.0583,
|
||||
"total": 15862,
|
||||
"unresolvedFailureCount": 8262
|
||||
},
|
||||
"warnings": [
|
||||
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
@@ -22016,13 +22019,12 @@
|
||||
"组合 Cambricon_mlu-370-x4|unknown|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b4|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_bi-150|transformers|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x8|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_mrv-100|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x8|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
]
|
||||
},
|
||||
"storageMode": "decision_state_only",
|
||||
"summarizedRecords": 15964,
|
||||
"summarizedRecords": 15967,
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -51,6 +51,7 @@
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:08:06.464776+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-bf16", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2340697867, "estimatedRequiredGiB": 2.621, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 2345554722, "modelscopeLicense": "other", "modelscopeParams": 1170340608, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2345554722}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.837081+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970332", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:08:06.464814+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248516007, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248516007}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.834872+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970331", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:08:06.464803+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Instruct-MLX-bf16", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2340697867, "estimatedRequiredGiB": 2.622, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 2345693224, "modelscopeLicense": "other", "modelscopeParams": 1170340608, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2345693224}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:37.819101+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970330", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-20T14:14:39.678925+00:00", "modelId": "aisingapore/SEA-LION-v1-7B", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "0335c7e4668aaa963303ebe91f80d6eb1e7e8d0299011bd85f7e05a4df031dc7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008115847, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008115847}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.801267+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970119", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["zaya"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:01:00.671056+00:00", "modelId": "JANGQ-AI/Osaurus-AppleScript-8B-JANG_6M", "modelProfile": {"architectures": ["ZayaForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7915599888, "estimatedRequiredGiB": 8.885, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "zaya", "modelscopeFileSize": 7950463197, "modelscopeLicense": "other", "modelscopeParams": 2274126328, "modelscopeTags": ["license:other", "model_type:zaya", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:applescript", "custom_tag:macos", "custom_tag:computer-use", "custom_tag:agent", "custom_tag:tool-calling", "custom_tag:function-calling", "custom_tag:mlx", "custom_tag:moe", "custom_tag:osaurus"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7950463197}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.783924+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970117", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:01:00.670830+00:00", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21930054875, "estimatedRequiredGiB": 24.532, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 21950415614, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6024588416, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:optiq", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21950415614}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:29.635942+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970115", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "transformers", "lastSyncTime": "2026-09-20T12:01:00.670989+00:00", "modelId": "aisingapore/SEA-LION-v1-7B-IT", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "0335c7e4668aaa963303ebe91f80d6eb1e7e8d0299011bd85f7e05a4df031dc7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361928, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008164431, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008164431}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:41:49.365180+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970023", "taskType": "text-generation", "verifyResult": -1}
|
||||
@@ -297,4 +298,3 @@
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["mellum"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-13T21:09:24.912761+00:00", "modelId": "JetBrains/Mellum2-12B-A2.5B-Instruct", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-13T21:09:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4152538", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm", "lastSyncTime": "2026-09-13T21:09:24.912775+00:00", "modelId": "mlx-community/Qwen3.5-122B-A10B-OptiQ-2bit-REAP-63B", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-13T21:05:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4490277", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-13T21:09:24.912708+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-13T21:01:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4336358", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "platform_infrastructure", "failureAction": "open_short_gpu_framework_circuit_and_retry_other_models", "failureCategory": "platform_infrastructure", "failureClassificationReason": "no_idle_device", "failureCode": null, "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "gpu_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-13T21:09:24.912745+00:00", "modelId": "BAAI/AquilaMed-RL", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-13T21:01:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4457974", "taskType": "text-generation", "verifyResult": -1}
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
File diff suppressed because it is too large
Load Diff
@@ -1,24 +1,24 @@
|
||||
{
|
||||
"agentVersion": "2026.09.20.2",
|
||||
"checksums": {
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "21c96a211d994bb7802d9b44acc2d44f9403f061ef7f28f13072aefa070d816f",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "992b389188f1dca7d2fa6160f3274d20ff5dd1c2e8551779e9a6d1c365545499",
|
||||
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
||||
".modelhub_state/market_intelligence.json": "981cbc0057e92007759bb7b7b2365a9fe6a8e1515a0bdfee65d05dd30641a3d1",
|
||||
".modelhub_state/official_capabilities.json": "6e28c2b5aca4f7512a0b238960415387861b5c01ceaeb02115b1f2220d999dda",
|
||||
".modelhub_state/outcome_checkpoint.json": "12db541eb571eec8e49248214ddd56e66951567154d4ea21fc36412c4345e6e7",
|
||||
".modelhub_state/market_intelligence.json": "64525c3613df5f8340dae6c81094ed0e38d27f3b7dd7fddca4b2b8b929607b1f",
|
||||
".modelhub_state/official_capabilities.json": "0a383113525b1fba3359ee6ef7ead1a6a4075184adcdb3fe32b67ac9e8f31cc0",
|
||||
".modelhub_state/outcome_checkpoint.json": "f1726fa082f227b3030a2656464c04156814bb9c64a6b0e5fc711ccee5f4e39f",
|
||||
".modelhub_state/queue_cleanup_latest.json": "eb01f106d0cd4018a0346c4a81182d6a9e67f5d4d14db22f5d0c685667977bc3",
|
||||
".modelhub_state/recent_outcomes.jsonl": "6973f12c807aad13e9919d00690216cc961734b6c02c5cb777c2686bd7109d80",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "3b32bcbabe4ef85c83e360d6443bdcd8223cda9af5410d620e535836477d00f5",
|
||||
".modelhub_state/recovery_intents.jsonl": "ba451c2abafb4936bfa615f29768031c194b82e87c1e54c989484570c585b1b8",
|
||||
".modelhub_state/recent_outcomes.jsonl": "306abefa1d94764ce0350dd8b4815acc6ec4602b6d992e6626fe9440d199eebb",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "1dc5ca356c786d9a2dd8258cbea673a28c92d49d05ed52b2e65736c79625bdce",
|
||||
".modelhub_state/recovery_intents.jsonl": "5580d6e970098701e3c6bf3ee7e7eb81088ac2b9444860c412494b63cc4d5df5",
|
||||
".modelhub_state/routing_intelligence.json": "118dae164974fbee630f675d1810a874d8a3847f8fc8d3df4b9ea0e460e45f4a",
|
||||
".modelhub_state/submission_exclusions.jsonl": "221be895a524330815b3c5124450b33c94b5f1e68169030c73684bf675d06ec7",
|
||||
".modelhub_state/worker_crashes.jsonl": "a494412048eca6dce83e7284303a6b4b56a445c1274ac201ab8723c739a14f23",
|
||||
"ledger/submissions.jsonl": "ec4279c2ddeba3d2b67abee40458baf34fc1c2c77a6b7db9ee7890ebb6627868",
|
||||
"outcomes/submissions.jsonl": "442647d17584d74946d5c2d75e8179b686678bfa02c7372967a0e67b218a2959"
|
||||
"outcomes/submissions.jsonl": "028356caa29db375b68b280713233e87803e11a617721930266199e42f4fc313"
|
||||
},
|
||||
"generation": 10612,
|
||||
"generation": 10613,
|
||||
"phase": "cycle",
|
||||
"schemaVersion": 1,
|
||||
"updatedAt": "2026-09-20T14:13:36.619136+00:00",
|
||||
"updatedAt": "2026-09-20T14:14:47.649255+00:00",
|
||||
"writerId": "fc715c06e24c4db2ae3e425e435db996"
|
||||
}
|
||||
|
||||
@@ -132,8 +132,6 @@
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-09T19:37:31.704706+00:00", "modelId": "OpenBMB/MiniCPM5-2B-Midtrain", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557128, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043777882, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043777882}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T11:33:17.847799+00:00", "targetGpu": "Vastai_va16", "taskId": "4739773", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-09T19:37:31.704762+00:00", "modelId": "OpenBMB/MiniCPM5-2B-Base", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557128, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043777904, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043777904}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T11:33:17.843220+00:00", "targetGpu": "Vastai_va16", "taskId": "4739774", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T19:56:42.896518+00:00", "modelId": "OpenBMB/MiniCPM5-2B-GPTQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2099615880, "estimatedRequiredGiB": 2.358, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 2109838556, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 2109838556}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T11:50:46.742747+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4739958", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T20:23:37.101884+00:00", "modelId": "OpenBMB/MiniCPM5-2B-Base", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557128, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043777904, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043777904}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T12:18:42.029908+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4740277", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-09T20:23:37.101831+00:00", "modelId": "OpenBMB/MiniCPM5-2B-Midtrain", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557128, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043777882, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043777882}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T12:18:42.031229+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4740276", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-09T23:05:47.802256+00:00", "modelId": "OpenBMB/MiniCPM5-2B-Midtrain", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557128, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043778298, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043778298}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T15:04:57.782677+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4742257", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-09T23:05:47.802235+00:00", "modelId": "OpenBMB/MiniCPM5-2B-Base", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557128, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043778320, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043778320}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T15:04:57.783545+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4742256", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-10T00:59:29.336134+00:00", "modelId": "laion/swesmith-nl2bash-stack-bugsseq", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8593, "estimatedRequiredGiB": 18.327, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 16397514624, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8190735360, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama-factory", "custom_tag:full", "custom_tag:generated_from_trainer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16398873171}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-09T16:58:40.742484+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4743575", "taskType": "text-generation", "verifyResult": null}
|
||||
@@ -842,7 +840,6 @@
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:01:00.670584+00:00", "modelId": "neuralmagic/gemma-2-2b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4385544296, "estimatedRequiredGiB": 4.926, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 4407367987, "modelscopeLicense": "gemma", "modelscopeParams": 3204165888, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4407367987}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:46:29.599298+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970107", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29274355224, "estimatedRequiredGiB": 32.761, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 29314389092, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 29314389092}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T03:46:35.789758+00:00", "targetGpu": "Biren_166m", "taskId": "4970118", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.670578+00:00", "modelId": "IntervitensInc/kek_mk3", "modelProfile": {"architectures": ["StableLmForCausalLM"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3289069288, "estimatedRequiredGiB": 3.683, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "stablelm", "modelscopeFileSize": 3295875324, "modelscopeLicense": null, "modelscopeParams": 1644515328, "modelscopeTags": ["model_type:stablelm", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3295875324}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:46:35.813804+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970120", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "transformers", "lastSyncTime": null, "modelId": "aisingapore/SEA-LION-v1-7B", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "0335c7e4668aaa963303ebe91f80d6eb1e7e8d0299011bd85f7e05a4df031dc7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008115847, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008115847}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-20T03:46:35.801267+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970119", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.670653+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B", "modelProfile": {"architectures": ["Lfm2MoeForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16936006912, "estimatedRequiredGiB": 18.948, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2_moe", "modelscopeFileSize": 16953948786, "modelscopeLicense": "other", "modelscopeParams": 8467856832, "modelscopeTags": ["license:other", "model_type:lfm2_moe", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16953948786}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:53:33.672925+00:00", "targetGpu": "Vastai_va16", "taskId": "4970243", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.670817+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663410366, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663410366}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:53:33.649508+00:00", "targetGpu": "Vastai_va16", "taskId": "4970244", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.670984+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:53:33.591582+00:00", "targetGpu": "Vastai_va16", "taskId": "4970241", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user