state: generation 5992 (cycle)

This commit is contained in:
2026-09-12 02:59:05 +00:00
parent da74f4fe1b
commit 54d1146595
11 changed files with 2176 additions and 2092 deletions

View File

@@ -1007,6 +1007,25 @@
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation"
},
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:nemotron_h": {
"architectureSignature": "model_type:nemotron_h",
"architectures": [],
"evidenceCount": 1,
"expiresAt": "2026-10-12T02:51:21+00:00",
"framework": "vllm_fix_tokenizer",
"latestFailureAt": "2026-09-12T02:51:21+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "nemotron_h",
"sourceModelIds": [
"nv-community/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-NVFP4"
],
"sourceTaskIds": [
"4134178"
],
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation"
},
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:qwen3_5": {
"architectureSignature": "model_type:qwen3_5",
"architectures": [],
@@ -1326,9 +1345,9 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-12T02:55:53.615460+00:00",
"generatedAt": "2026-09-12T02:59:04.590587+00:00",
"summary": {
"activeBlockCount": 67,
"activeBlockCount": 68,
"byGpuFramework": {
"Ascend_910-b3|vllm": 8,
"Ascend_910-b4|vllm": 9,
@@ -1338,7 +1357,7 @@
"MetaX_c-500|vllm": 7,
"Mthreads_s4000|vllm": 7,
"Sunrise_pt-200-x1|vllm": 1,
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 2,
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 3,
"Vastai_va16|vllm": 12,
"Vastai_va16|vllm_fix_tokenizer": 2,
"hygon_k100-ai|vllm": 7,

View File

@@ -35,7 +35,7 @@
"2": {
"complete": false,
"lastError": "ModelHubAPIError: 系统错误",
"listingErrors": 430,
"listingErrors": 431,
"nextPage": 1,
"recordsScanned": 0,
"uniqueRecords": 0
@@ -102,12 +102,12 @@
"cutoffAt": "2026-09-04T03:55:51.365685+00:00",
"failureLogsInspected": 0,
"mode": "incremental_decision_only",
"nextAccountIndex": 2,
"nextAccountIndex": 3,
"recordsScanned": 0,
"seenTaskIds": [],
"startedAt": "2026-09-04T03:55:51.365685+00:00",
"terminalRecords": 0,
"uniqueRecords": 0,
"updatedAt": "2026-09-12T02:55:53.587554+00:00",
"updatedAt": "2026-09-12T02:59:04.560993+00:00",
"version": 1
}

View File

@@ -416,7 +416,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-12T02:54:17.744485+00:00",
"generatedAt": "2026-09-12T02:57:05.481546+00:00",
"gpuStats": {
"Ascend_910-b3": {
"available": true,

File diff suppressed because it is too large Load Diff

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-12T02:48:11.405538+00:00",
"lastSyncTime": "2026-09-12T02:48:11.158163+00:00",
"generatedAt": "2026-09-12T02:56:54.666019+00:00",
"lastSyncTime": "2026-09-12T02:56:54.384007+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -1011,6 +1011,25 @@
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation"
},
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:nemotron_h": {
"architectureSignature": "model_type:nemotron_h",
"architectures": [],
"evidenceCount": 1,
"expiresAt": "2026-10-12T02:51:21+00:00",
"framework": "vllm_fix_tokenizer",
"latestFailureAt": "2026-09-12T02:51:21+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "nemotron_h",
"sourceModelIds": [
"nv-community/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-NVFP4"
],
"sourceTaskIds": [
"4134178"
],
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation"
},
"sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation|model_type:qwen3_5": {
"architectureSignature": "model_type:qwen3_5",
"architectures": [],
@@ -1331,7 +1350,7 @@
}
},
"architectureCompatibilitySummary": {
"activeBlockCount": 67,
"activeBlockCount": 68,
"byGpuFramework": {
"Ascend_910-b3|vllm": 8,
"Ascend_910-b4|vllm": 9,
@@ -1341,7 +1360,7 @@
"MetaX_c-500|vllm": 7,
"Mthreads_s4000|vllm": 7,
"Sunrise_pt-200-x1|vllm": 1,
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 2,
"Sunrise_pt-200-x1|vllm_fix_tokenizer": 3,
"Vastai_va16|vllm": 12,
"Vastai_va16|vllm_fix_tokenizer": 2,
"hygon_k100-ai|vllm": 7,
@@ -1482,18 +1501,18 @@
"unresolvedFailureCount": 1
},
"Ascend_910-b4|vllm|text-generation": {
"attributableFailureCount": 37,
"attributableFailureCount": 38,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 37,
"decisionTotal": 38,
"failureBreakdown": {
"ambiguous_runtime": 33,
"framework_architecture_unsupported": 33,
"framework_architecture_unsupported": 34,
"memory_capacity": 1,
"repository_structure": 1,
"tokenizer_compatibility": 2
},
"failureCount": 70,
"failureCount": 71,
"failureRate": 1.0,
"framework": "vllm",
"pendingCount": 0,
@@ -1503,7 +1522,7 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b4",
"taskType": "text-generation",
"total": 70,
"total": 71,
"unresolvedFailureCount": 33
},
"Biren_166m|unknown|text-generation": {
@@ -2082,17 +2101,17 @@
"unresolvedFailureCount": 33
},
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": {
"attributableFailureCount": 41,
"attributableFailureCount": 42,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 41,
"decisionTotal": 42,
"failureBreakdown": {
"framework_architecture_unsupported": 3,
"framework_architecture_unsupported": 4,
"platform_infrastructure": 2,
"tokenizer_compatibility": 38,
"参数/模板问题": 6
},
"failureCount": 49,
"failureCount": 50,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"pendingCount": 0,
@@ -2102,22 +2121,22 @@
"successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation",
"total": 49,
"total": 50,
"unresolvedFailureCount": 6
},
"Sunrise_pt-200-x1|vllm|text-generation": {
"attributableFailureCount": 2,
"attributableFailureCount": 3,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 2,
"decisionTotal": 3,
"failureBreakdown": {
"ambiguous_runtime": 1,
"framework_architecture_unsupported": 1,
"platform_infrastructure": 1,
"tokenizer_compatibility": 1,
"tokenizer_compatibility": 2,
"参数/模板问题": 2
},
"failureCount": 6,
"failureCount": 7,
"failureRate": 1.0,
"framework": "vllm",
"pendingCount": 0,
@@ -2127,7 +2146,7 @@
"successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation",
"total": 6,
"total": 7,
"unresolvedFailureCount": 3
},
"Vastai_va16|unknown|text-generation": {
@@ -2356,30 +2375,30 @@
"unresolvedFailureCount": 862
},
"vllm": {
"attributableFailureCount": 246,
"attributableFailureCount": 248,
"decisionFailureRate": 0.988,
"decisionSuccessRate": 0.012,
"decisionTotal": 249,
"decisionTotal": 251,
"failureBreakdown": {
"ambiguous_runtime": 195,
"backend_operator": 33,
"framework_architecture_unsupported": 172,
"framework_architecture_unsupported": 173,
"memory_capacity": 6,
"model_load": 20,
"platform_infrastructure": 1,
"repository_structure": 5,
"runtime_memory": 6,
"tokenizer_compatibility": 4,
"tokenizer_compatibility": 5,
"参数/模板问题": 24
},
"failureCount": 466,
"failureCount": 468,
"failureRate": 0.9936,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 1,
"successCount": 3,
"successRate": 0.0064,
"total": 469,
"total": 471,
"unresolvedFailureCount": 219
},
"vllm-mlu": {
@@ -2442,26 +2461,26 @@
"unresolvedFailureCount": 4
},
"vllm_fix_tokenizer": {
"attributableFailureCount": 44,
"attributableFailureCount": 45,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 44,
"decisionTotal": 45,
"failureBreakdown": {
"ambiguous_runtime": 36,
"framework_architecture_unsupported": 5,
"framework_architecture_unsupported": 6,
"model_load": 1,
"platform_infrastructure": 2,
"tokenizer_compatibility": 38,
"参数/模板问题": 7
},
"failureCount": 89,
"failureCount": 90,
"failureRate": 1.0,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 0,
"successRate": 0.0,
"total": 89,
"total": 90,
"unresolvedFailureCount": 43
},
"vllm_tokenizer_patch": {
@@ -2483,7 +2502,7 @@
"unresolvedFailureCount": 1
}
},
"generatedAt": "2026-09-12T02:48:11.401315+00:00",
"generatedAt": "2026-09-12T02:56:54.660525+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 26,
@@ -2509,13 +2528,13 @@
"unresolvedFailureCount": 63
},
"Ascend_910-b4": {
"attributableFailureCount": 38,
"decisionFailureRate": 0.95,
"decisionSuccessRate": 0.05,
"decisionTotal": 40,
"attributableFailureCount": 39,
"decisionFailureRate": 0.9512,
"decisionSuccessRate": 0.0488,
"decisionTotal": 41,
"failureBreakdown": {
"ambiguous_runtime": 34,
"framework_architecture_unsupported": 33,
"framework_architecture_unsupported": 34,
"memory_capacity": 1,
"model_load": 1,
"repository_structure": 1,
@@ -2523,14 +2542,14 @@
"参数/模板问题": 12,
"验证失败": 175
},
"failureCount": 259,
"failureRate": 0.9923,
"failureCount": 260,
"failureRate": 0.9924,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 2,
"successRate": 0.0077,
"total": 261,
"successRate": 0.0076,
"total": 262,
"unresolvedFailureCount": 221
},
"Biren_166m": {
@@ -2741,26 +2760,26 @@
"unresolvedFailureCount": 77
},
"Sunrise_pt-200-x1": {
"attributableFailureCount": 43,
"decisionFailureRate": 0.9773,
"decisionSuccessRate": 0.0227,
"decisionTotal": 44,
"attributableFailureCount": 45,
"decisionFailureRate": 0.9783,
"decisionSuccessRate": 0.0217,
"decisionTotal": 46,
"failureBreakdown": {
"ambiguous_runtime": 1,
"framework_architecture_unsupported": 4,
"framework_architecture_unsupported": 5,
"platform_infrastructure": 3,
"tokenizer_compatibility": 39,
"tokenizer_compatibility": 40,
"参数/模板问题": 9,
"验证失败": 32
},
"failureCount": 88,
"failureRate": 0.9888,
"failureCount": 90,
"failureRate": 0.989,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 3,
"successCount": 1,
"successRate": 0.0112,
"total": 89,
"successRate": 0.011,
"total": 91,
"unresolvedFailureCount": 42
},
"Vastai_va16": {
@@ -3976,6 +3995,29 @@
"total": 1,
"unresolvedFailureCount": 0
},
"Sunrise_pt-200-x1|vllm|text-generation|llama|compressed-tensors": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"tokenizer_compatibility": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"modelType": "llama",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Vastai_va16|vllm_fix_tokenizer|text-generation|hrm_text|none": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
@@ -4204,7 +4246,7 @@
"failureRate": 1.0,
"framework": "vllm",
"lastPlatformFailureAt": null,
"lastTerminalAt": "2026-09-10T15:48:17.745912+00:00",
"lastTerminalAt": "2026-09-12T02:56:54.384007+00:00",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
@@ -4551,22 +4593,22 @@
"unresolvedFailureCount": 0
},
"Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation": {
"attributableFailureCount": 18,
"consecutiveFailures": 18,
"attributableFailureCount": 17,
"consecutiveFailures": 17,
"consecutivePlatformFailures": 0,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 18,
"decisionTotal": 17,
"failureBreakdown": {
"framework_architecture_unsupported": 2,
"platform_infrastructure": 2,
"tokenizer_compatibility": 16
"tokenizer_compatibility": 15
},
"failureCount": 20,
"failureCount": 19,
"failureRate": 1.0,
"framework": "vllm_fix_tokenizer",
"lastPlatformFailureAt": "2026-09-10T11:01:21+00:00",
"lastTerminalAt": "2026-09-12T02:48:11.158114+00:00",
"lastTerminalAt": "2026-09-12T02:56:54.383989+00:00",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
@@ -4574,7 +4616,7 @@
"successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation",
"total": 20,
"total": 19,
"unresolvedFailureCount": 0
},
"Sunrise_pt-200-x1|vllm|text-generation": {
@@ -6928,6 +6970,30 @@
"total": 1,
"unresolvedFailureCount": 0
},
"Sunrise_pt-200-x1|vllm|text-generation|llama|compressed-tensors|30": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"tokenizer_compatibility": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm",
"loadSizeLog2Bucket": 30,
"modelType": "llama",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Sunrise_pt-200-x1",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Vastai_va16|vllm_fix_tokenizer|text-generation|hrm_text|none|31": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
@@ -7121,41 +7187,40 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 1677,
"totalRecords": 1764,
"terminalRecords": 1680,
"totalRecords": 1767,
"totals": {
"attributableFailureCount": 449,
"decisionFailureRate": 0.8891,
"decisionSuccessRate": 0.1109,
"decisionTotal": 505,
"attributableFailureCount": 452,
"decisionFailureRate": 0.8898,
"decisionSuccessRate": 0.1102,
"decisionTotal": 508,
"failureBreakdown": {
"ambiguous_runtime": 394,
"backend_operator": 40,
"framework_architecture_unsupported": 278,
"framework_architecture_unsupported": 280,
"memory_capacity": 11,
"model_load": 64,
"platform_infrastructure": 4,
"repository_structure": 5,
"runtime_memory": 6,
"tokenizer_compatibility": 45,
"tokenizer_compatibility": 46,
"参数/模板问题": 101,
"验证失败": 673
},
"failureCount": 1621,
"failureRate": 0.9666,
"failureCount": 1624,
"failureRate": 0.9667,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 4,
"successCount": 56,
"successRate": 0.0334,
"total": 1677,
"successRate": 0.0333,
"total": 1680,
"unresolvedFailureCount": 1168
},
"warnings": [
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Sunrise_pt-200-x1 本地统计失败率偏高≥50%),建议重点关注。",
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
@@ -7164,6 +7229,7 @@
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Sunrise_pt-200-x1 本地统计失败率偏高≥50%),建议重点关注。",
"组合 hygon_k100-ai|vllm-patch-tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -7179,11 +7245,12 @@
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Vastai_va16|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 1764,
"summarizedRecords": 1767,
"version": 1
}

View File

@@ -14,13 +14,13 @@
100
],
"accounts": 12,
"activeScanned": 1055,
"activeScanned": 1054,
"ageCleanupMode": "admission_only",
"agePolicySkipped": {
"cleanupDisabled": true,
"reason": "admission_only"
},
"architectureBlockCount": 67,
"architectureBlockCount": 68,
"architectureFrameworkCatalog": {
"ascend_910-b3|text-generation": [
"llamacpp",
@@ -98,10 +98,10 @@
"architectureOnly": true,
"architecturePolicySkipped": {
"frameworkCatalogUnknown": 87,
"frameworkContextUnknown": 330,
"frameworkContextUnknown": 329,
"modelArchitectureUnknown": 150,
"noMatchingBlock": 720,
"partiallyBlockedFrameworkSet": 98,
"noMatchingBlock": 718,
"partiallyBlockedFrameworkSet": 99,
"runningMatchedProtected": 0,
"submissionContextMismatch": 0,
"submissionContextUnknown": 0

View File

@@ -1,3 +1,5 @@
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-12T02:56:54.384007+00:00", "modelId": "amd/Qwen2.5-7B-Instruct-onnx-ryzenai-hybrid", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:53:23+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4079108", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["nemotron_h"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:56:54.383989+00:00", "modelId": "nv-community/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-NVFP4", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:51:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4134178", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:48:11.158114+00:00", "modelId": "RedHatAI/SmolLM-360M-Instruct-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:43:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4458009", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:39:29.219599+00:00", "modelId": "RedHatAI/Qwen2.5-3B-quantized.w4a16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:39:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4100876", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["mellum"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-12T02:39:29.219629+00:00", "modelId": "JetBrains/Mellum2-12B-A2.5B-Thinking", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-12T02:35:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4152611", "taskType": "text-generation", "verifyResult": -1}
@@ -296,5 +298,3 @@
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T13:16:14.206275+00:00", "modelId": "neuralmagic/gemma-2-9b-it-quantized.w4a16", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:57:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4107036", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:42:19.103212+00:00", "modelId": "neuralmagic/Qwen2-0.5B-Instruct-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:21:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4133653", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:42:19.103221+00:00", "modelId": "RedHatAI/Llama-2-7b-chat-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:19:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4107014", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:42:19.103172+00:00", "modelId": "iic/UEmbed-4B", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:15:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4152744", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "tokenizer_compatibility", "failureAction": "check_tokenizer_files_then_llm", "failureCategory": "tokenizer_compatibility", "failureClassificationReason": "broad_tokenizer_error", "failureCode": "TOKENIZER_FAILED", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_framework", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-08T12:11:42.795229+00:00", "modelId": "RedHatAI/starcoder2-3b-quantized.w8a8", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T12:11:21+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4101032", "taskType": "text-generation", "verifyResult": -1}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -1,24 +1,24 @@
{
"agentVersion": "2026.09.04.4",
"checksums": {
".modelhub_state/architecture_compatibility_blacklist.json": "011eb3c1b849bf9708f1bac75ca4324fba7ccea5b2d31e7958672b438ac07a2e",
".modelhub_state/architecture_history_backfill.json": "c40825e111c98ca7ec5253810d8b6316b151ae07f452172b53490b317ef8b8f9",
".modelhub_state/market_intelligence.json": "425fd2d055e24890c3dbdcf7d3c53fa3b5d026de8bcd0ea2e575042233b6db62",
".modelhub_state/official_capabilities.json": "8369b20d67b9198cfbc98c63f085f155c2ad542fce1d7c0bb89011603d855dc4",
".modelhub_state/outcome_checkpoint.json": "267592fed01a3ba5fb68028d166d6fe5c608ecec7c2b32b1985bc4f6099f8d69",
".modelhub_state/queue_cleanup_latest.json": "384aaa478516df0fbe81d4231e5cdb026ef79a424962096f691840bfd3cb6931",
".modelhub_state/recent_outcomes.jsonl": "f40a93b9569ca8a0358c678853f3510722522feb6de5c702465195362b816654",
".modelhub_state/recovery_active_tasks.jsonl": "809d9b049c3245950f42f45c35800399af62b9854ae42055db6811592158e5b2",
".modelhub_state/recovery_intents.jsonl": "74ebea4b0eafef707cdb83c55ad4ad7b0ac81b3545e148e9bdbf592a944c91cd",
".modelhub_state/architecture_compatibility_blacklist.json": "721e1bad1444c6c604a303eb5d9f50044d3c507e81060f0058923b4e695c5038",
".modelhub_state/architecture_history_backfill.json": "3f2e1bd8a9f7ed99a909f2d1254ac7953d404478b10472da6b49167fd2932efc",
".modelhub_state/market_intelligence.json": "4f0c21f848715ec0eaac426844286809622bcc6e62d0249a2139bf7a621b61e2",
".modelhub_state/official_capabilities.json": "158704e2ce643a87d7843a5369834db9339a6c2db1c7e976c5b86eb55039ce75",
".modelhub_state/outcome_checkpoint.json": "de8419898d771b485c1aeaedd1aa952b409d96ddd876a51daf2d8e645433c46f",
".modelhub_state/queue_cleanup_latest.json": "dbff8ecbdcba815a2302661742ce4689464c0250a7ef3678d448f0d7a5dc07da",
".modelhub_state/recent_outcomes.jsonl": "fea2c8180e648ab702cfb751b88efcae2b55820fdeb03576bab6e5108dbdc0a3",
".modelhub_state/recovery_active_tasks.jsonl": "dd85494be676f6117957f3ec4e84bf7b558df9553e5945549601a22113cef06b",
".modelhub_state/recovery_intents.jsonl": "cebd8ca56df5eb250c64975557e8591f5e20a5e8483360ade4410f572a5434e9",
".modelhub_state/routing_intelligence.json": "316b7d02e32cf2dd3f7c0ca452f58d571a371863ac4d1c1b5193456035ce3a53",
".modelhub_state/submission_exclusions.jsonl": "14480c58228be6b76e21e98008960b815d20d16132dcab7a11e5b449c5e2d220",
".modelhub_state/worker_crashes.jsonl": "9693a31a13cc18a3136ff5373f9569dc2fecaa274aaf91d3b1c8a67f31fe0162",
"ledger/submissions.jsonl": "8abcf838f382e71b9140db2901b93972a33b043c4c67a56473d00af9ccd9be3c",
"outcomes/submissions.jsonl": "c53f6ac33f79e616e13d79bc0ead6c3f607869ad8325b1aec0fae42c5c18de72"
"outcomes/submissions.jsonl": "89e4ed1669c9c8d079363347a986e7e020d5688c4a0a5b51c8c2e756ca106cb7"
},
"generation": 5991,
"generation": 5992,
"phase": "cycle",
"schemaVersion": 1,
"updatedAt": "2026-09-12T02:55:53.702813+00:00",
"updatedAt": "2026-09-12T02:59:05.238590+00:00",
"writerId": "e5dd59a4c84a4f21926bad61448789e2"
}

View File

@@ -81,7 +81,6 @@
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-05T07:14:11.020202+00:00", "modelId": "inferencerlabs/DeepSeek-V4-Flash-MTP-DSpark-MLX", "modelProfile": {"architectures": [], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12994906286, "estimatedRequiredGiB": 14.53, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "deepseek_v4_dspark", "modelscopeFileSize": 13001289422, "modelscopeLicense": null, "modelscopeParams": 4276397927, "modelscopeTags": ["model_type:deepseek_v4_dspark", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13001289422}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:12:12.056590+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4636059", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T07:14:11.020212+00:00", "modelId": "neuralmagic/SmolLM-360M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ac9457423139c55bd60f70bc9092d1d9dddd724b01ac0d5cac27be6ed145e50e", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 504052632, "estimatedRequiredGiB": 0.567, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 507443106, "modelscopeLicense": "apache-2.0", "modelscopeParams": 409007040, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 507443106}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:12:18.663255+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4636060", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T07:14:11.020096+00:00", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5705684944, "estimatedRequiredGiB": 6.387, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5714910031, "modelscopeLicense": null, "modelscopeParams": 8031506432, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 5714910031}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:12:29.786627+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4636099", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T07:21:16.625484+00:00", "modelId": "nm-testing/tinyllama-oneshot-w8a8-static-v2", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "ac9457423139c55bd60f70bc9092d1d9dddd724b01ac0d5cac27be6ed145e50e", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1231270112, "estimatedRequiredGiB": 1.378, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 1233118759, "modelscopeLicense": null, "modelscopeParams": 1100048384, "modelscopeTags": ["model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 1233118759}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:20:09.476751+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4636203", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-05T07:32:32.823940+00:00", "modelId": "XHToken/Spark-X2.5-1.7B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3415340496, "estimatedRequiredGiB": 3.834, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://cambricon.com/index.php?a=lists&c=index&catid=406&m=content"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 3430847031, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1707657216, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3430847031}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:29:59.902384+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4636365", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-05T07:32:32.823946+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-v0.4-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727938576, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://cambricon.com/index.php?a=lists&c=index&catid=406&m=content"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737300809, "modelscopeLicense": "other", "modelscopeParams": 8030261248, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737300809}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:29:59.915927+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4636369", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-05T07:32:32.823927+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-DPO-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://cambricon.com/index.php?a=lists&c=index&catid=406&m=content"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737317601, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737317601}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T23:29:59.925332+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4636370", "taskType": "text-generation", "verifyResult": null}
@@ -618,8 +617,8 @@
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T01:25:57.322160+00:00", "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T17:22:24.744587+00:00", "targetGpu": "MetaX_c-500", "taskId": "4783634", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-12T02:21:54.921668+00:00", "modelId": "guaidao2/LFM2.5-2.6B-For-CTF", "modelProfile": {"architectures": [], "configFingerprint": "6647f6920340e837cb19ab6ec8bc02c4d5d8cebe4d720e359319e40aa1ebaf79", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2874779040, "estimatedRequiredGiB": 11.734, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 10499570876, "modelscopeLicense": null, "modelscopeParams": 2697198592, "modelscopeTags": ["library:gguf", "task:text-generation", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 10499570876}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:18:29.685555+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4784429", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T02:21:54.921691+00:00", "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:18:29.687478+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4784430", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T18:51:22.774816+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784981", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12999919368, "estimatedRequiredGiB": 14.555, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 13023155398, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3544073216, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:imatrix", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13023155398}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T18:51:22.770505+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784982", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T02:56:54.383907+00:00", "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:51:22.774816+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784981", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-12T02:56:54.383964+00:00", "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12999919368, "estimatedRequiredGiB": 14.555, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 13023155398, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3544073216, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:imatrix", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13023155398}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-11T18:51:22.770505+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4784982", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730839322, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-12T01:22:51.902664+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4791408", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12999919368, "estimatedRequiredGiB": 14.555, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 13023155398, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3544073216, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:imatrix", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 2}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-10T18:58:50.012615+00:00", "modelType": "qwen3_5", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 2}, "repositoryOnDiskBytes": 13023155398}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-12T01:22:51.904770+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4791409", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "nex-agi/Nex-N2.5-mini", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 70214494016, "estimatedRequiredGiB": 78.495, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 70235999301, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35107181936, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 2}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-09T14:00:25.508497+00:00", "modelType": "qwen3_5_moe", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 2}, "repositoryOnDiskBytes": 70235999301}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-12T01:22:51.908714+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4791410", "taskType": "text-generation", "verifyResult": null}