state: generation 18477 (cycle)

This commit is contained in:
2026-09-29 17:58:00 +00:00
parent 2c601b6b35
commit dd89cc7a2d
10 changed files with 1660 additions and 1578 deletions

View File

@@ -786,6 +786,25 @@
"targetGpu": "Biren_166m",
"taskType": "text-generation"
},
"cambricon_mlu-370-x4|vllm-customized|text-generation|model_type:qwen3_5": {
"architectureSignature": "model_type:qwen3_5",
"architectures": [],
"evidenceCount": 1,
"expiresAt": "2026-10-22T05:27:11.098010+00:00",
"framework": "vllm-customized",
"latestFailureAt": "2026-09-22T05:27:11.098010+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "qwen3_5",
"sourceModelIds": [
"Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed"
],
"sourceTaskIds": [
"5012909"
],
"targetGpu": "Cambricon_mlu-370-x4",
"taskType": "text-generation"
},
"cambricon_mlu-370-x4|vllm|text-generation|model_type:lfm2": {
"architectureSignature": "model_type:lfm2",
"architectures": [],
@@ -3280,9 +3299,9 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-29T17:45:52.762781+00:00",
"generatedAt": "2026-09-29T17:54:53.642992+00:00",
"summary": {
"activeBlockCount": 165,
"activeBlockCount": 166,
"byGpuFramework": {
"Ascend_910-b3|vllm": 13,
"Ascend_910-b3|vllm_tokenizer_patch": 5,
@@ -3290,6 +3309,7 @@
"Ascend_910-b4|vllm_tokenizer_patch": 6,
"Biren_166m|vllm": 3,
"Cambricon_mlu-370-x4|vllm": 4,
"Cambricon_mlu-370-x4|vllm-customized": 1,
"Cambricon_mlu-370-x8|vllm": 16,
"Cambricon_mlu-370-x8|vllm-customized": 6,
"Cambricon_mlu-370-x8|vllm-mlu": 5,

View File

@@ -434,7 +434,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-29T17:53:50.865820+00:00",
"generatedAt": "2026-09-29T17:57:05.751644+00:00",
"gpuStats": {
"Ascend_910-b3": {
"available": true,

View File

@@ -1,5 +1,5 @@
{
"catalogUpdatedAt": "2026-09-29T17:53:50.865820+00:00",
"catalogUpdatedAt": "2026-09-29T17:57:05.751644+00:00",
"configuredTaskTypes": [
"text-generation"
],
@@ -56,7 +56,7 @@
"time-series-forecasting"
],
"errors": [],
"generatedAt": "2026-09-29T17:53:50.865820+00:00",
"generatedAt": "2026-09-29T17:57:05.751644+00:00",
"gpuCatalog": {
"Ascend_910-b3": {
"canVerify": true,
@@ -6316,6 +6316,6 @@
"updateTime": "2025-12-22 08:59:53"
}
],
"taskTreeUpdatedAt": "2026-09-29T17:53:50.865820+00:00",
"taskTreeUpdatedAt": "2026-09-29T17:57:05.751644+00:00",
"version": 1
}

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-29T17:16:53.405682+00:00",
"lastSyncTime": "2026-09-29T17:16:53.043636+00:00",
"generatedAt": "2026-09-29T17:54:53.540400+00:00",
"lastSyncTime": "2026-09-29T17:54:53.270754+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -790,6 +790,25 @@
"targetGpu": "Biren_166m",
"taskType": "text-generation"
},
"cambricon_mlu-370-x4|vllm-customized|text-generation|model_type:qwen3_5": {
"architectureSignature": "model_type:qwen3_5",
"architectures": [],
"evidenceCount": 1,
"expiresAt": "2026-10-22T05:27:11.098010+00:00",
"framework": "vllm-customized",
"latestFailureAt": "2026-09-22T05:27:11.098010+00:00",
"latestSuccessfulAt": null,
"matchType": "model_type",
"modelType": "qwen3_5",
"sourceModelIds": [
"Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed"
],
"sourceTaskIds": [
"5012909"
],
"targetGpu": "Cambricon_mlu-370-x4",
"taskType": "text-generation"
},
"cambricon_mlu-370-x4|vllm|text-generation|model_type:lfm2": {
"architectureSignature": "model_type:lfm2",
"architectures": [],
@@ -3285,7 +3304,7 @@
}
},
"architectureCompatibilitySummary": {
"activeBlockCount": 165,
"activeBlockCount": 166,
"byGpuFramework": {
"Ascend_910-b3|vllm": 13,
"Ascend_910-b3|vllm_tokenizer_patch": 5,
@@ -3293,6 +3312,7 @@
"Ascend_910-b4|vllm_tokenizer_patch": 6,
"Biren_166m|vllm": 3,
"Cambricon_mlu-370-x4|vllm": 4,
"Cambricon_mlu-370-x4|vllm-customized": 1,
"Cambricon_mlu-370-x8|vllm": 16,
"Cambricon_mlu-370-x8|vllm-customized": 6,
"Cambricon_mlu-370-x8|vllm-mlu": 5,
@@ -3402,28 +3422,28 @@
"unresolvedFailureCount": 1
},
"Ascend_910-b3|vllm_tokenizer_patch|text-generation": {
"attributableFailureCount": 34,
"decisionFailureRate": 0.8947,
"decisionSuccessRate": 0.1053,
"decisionTotal": 38,
"attributableFailureCount": 35,
"decisionFailureRate": 0.8974,
"decisionSuccessRate": 0.1026,
"decisionTotal": 39,
"failureBreakdown": {
"ambiguous_runtime": 64,
"context_length": 5,
"context_length": 6,
"framework_architecture_unsupported": 28,
"memory_capacity": 1,
"参数/模板问题": 8
},
"failureCount": 106,
"failureRate": 0.9636,
"failureCount": 107,
"failureRate": 0.964,
"framework": "vllm_tokenizer_patch",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 4,
"successRate": 0.0364,
"successRate": 0.036,
"targetGpu": "Ascend_910-b3",
"taskType": "text-generation",
"total": 110,
"total": 111,
"unresolvedFailureCount": 72
},
"Ascend_910-b3|vllm|text-generation": {
@@ -3837,15 +3857,15 @@
"unresolvedFailureCount": 5
},
"Cambricon_mlu-370-x4|vllm-customized|text-generation": {
"attributableFailureCount": 2,
"attributableFailureCount": 3,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 2,
"decisionTotal": 3,
"failureBreakdown": {
"ambiguous_runtime": 4,
"framework_architecture_unsupported": 2
"framework_architecture_unsupported": 3
},
"failureCount": 6,
"failureCount": 7,
"failureRate": 1.0,
"framework": "vllm-customized",
"pendingCount": 0,
@@ -3855,7 +3875,7 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x4",
"taskType": "text-generation",
"total": 6,
"total": 7,
"unresolvedFailureCount": 4
},
"Cambricon_mlu-370-x4|vllm-mlu|text-generation": {
@@ -6006,26 +6026,26 @@
"unresolvedFailureCount": 1637
},
"vllm-customized": {
"attributableFailureCount": 14,
"attributableFailureCount": 15,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 14,
"decisionTotal": 15,
"failureBreakdown": {
"ambiguous_runtime": 46,
"framework_architecture_unsupported": 10,
"framework_architecture_unsupported": 11,
"model_load": 3,
"platform_infrastructure": 2,
"tokenizer_compatibility": 1,
"参数/模板问题": 5
},
"failureCount": 67,
"failureCount": 68,
"failureRate": 1.0,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 0,
"successRate": 0.0,
"total": 67,
"total": 68,
"unresolvedFailureCount": 51
},
"vllm-mlu": {
@@ -6125,14 +6145,14 @@
"unresolvedFailureCount": 342
},
"vllm_tokenizer_patch": {
"attributableFailureCount": 131,
"decisionFailureRate": 0.9291,
"decisionSuccessRate": 0.0709,
"decisionTotal": 141,
"attributableFailureCount": 132,
"decisionFailureRate": 0.9296,
"decisionSuccessRate": 0.0704,
"decisionTotal": 142,
"failureBreakdown": {
"ambiguous_runtime": 176,
"backend_operator": 15,
"context_length": 5,
"context_length": 6,
"framework_architecture_unsupported": 78,
"memory_capacity": 2,
"model_load": 23,
@@ -6142,27 +6162,27 @@
"tokenizer_compatibility": 2,
"参数/模板问题": 29
},
"failureCount": 338,
"failureCount": 339,
"failureRate": 0.9713,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 10,
"successRate": 0.0287,
"total": 348,
"total": 349,
"unresolvedFailureCount": 205
}
},
"generatedAt": "2026-09-29T17:16:53.391898+00:00",
"generatedAt": "2026-09-29T17:54:53.527997+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 104,
"decisionFailureRate": 0.8254,
"decisionSuccessRate": 0.1746,
"decisionTotal": 126,
"attributableFailureCount": 105,
"decisionFailureRate": 0.8268,
"decisionSuccessRate": 0.1732,
"decisionTotal": 127,
"failureBreakdown": {
"ambiguous_runtime": 130,
"context_length": 5,
"context_length": 6,
"framework_architecture_unsupported": 93,
"memory_capacity": 2,
"repository_structure": 1,
@@ -6171,14 +6191,14 @@
"日志缺失": 3,
"验证失败": 27
},
"failureCount": 303,
"failureRate": 0.9323,
"failureCount": 304,
"failureRate": 0.9325,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 22,
"successRate": 0.0677,
"total": 325,
"successRate": 0.0675,
"total": 326,
"unresolvedFailureCount": 199
},
"Ascend_910-b4": {
@@ -6240,15 +6260,15 @@
"unresolvedFailureCount": 558
},
"Cambricon_mlu-370-x4": {
"attributableFailureCount": 748,
"decisionFailureRate": 0.9111,
"decisionSuccessRate": 0.0889,
"decisionTotal": 821,
"attributableFailureCount": 749,
"decisionFailureRate": 0.9112,
"decisionSuccessRate": 0.0888,
"decisionTotal": 822,
"failureBreakdown": {
"ambiguous_runtime": 508,
"architecture_compatibility": 53,
"context_length": 48,
"framework_architecture_unsupported": 301,
"framework_architecture_unsupported": 302,
"memory_capacity": 219,
"model_load": 7,
"platform_infrastructure": 16,
@@ -6258,14 +6278,14 @@
"日志缺失": 13,
"验证失败": 42
},
"failureCount": 1614,
"failureRate": 0.9567,
"failureCount": 1615,
"failureRate": 0.9568,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 16,
"successCount": 73,
"successRate": 0.0433,
"total": 1687,
"successRate": 0.0432,
"total": 1688,
"unresolvedFailureCount": 850
},
"Cambricon_mlu-370-x8": {
@@ -6845,15 +6865,15 @@
"unresolvedFailureCount": 6
},
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|llama|compressed-tensors": {
"attributableFailureCount": 3,
"attributableFailureCount": 4,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 3,
"decisionTotal": 4,
"failureBreakdown": {
"ambiguous_runtime": 11,
"context_length": 3
"context_length": 4
},
"failureCount": 14,
"failureCount": 15,
"failureRate": 1.0,
"framework": "vllm_tokenizer_patch",
"modelType": "llama",
@@ -6865,7 +6885,7 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b3",
"taskType": "text-generation",
"total": 14,
"total": 15,
"unresolvedFailureCount": 11
},
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|llama|none": {
@@ -10730,6 +10750,29 @@
"total": 1,
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x4|vllm-customized|text-generation|qwen3_5|none": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-customized",
"modelType": "qwen3_5",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x4",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Cambricon_mlu-370-x4|vllm-customized|text-generation|stablelm|none": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -23499,20 +23542,21 @@
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x4|vllm-customized|text-generation": {
"attributableFailureCount": 0,
"consecutiveFailures": 0,
"attributableFailureCount": 1,
"consecutiveFailures": 1,
"consecutivePlatformFailures": 0,
"decisionFailureRate": 0.0,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"decisionTotal": 1,
"failureBreakdown": {
"ambiguous_runtime": 1
"ambiguous_runtime": 1,
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureCount": 2,
"failureRate": 1.0,
"framework": "vllm-customized",
"lastPlatformFailureAt": null,
"lastTerminalAt": "2026-09-27T06:27:28.565882+00:00",
"lastTerminalAt": "2026-09-29T17:54:53.270719+00:00",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
@@ -23520,7 +23564,7 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x4",
"taskType": "text-generation",
"total": 1,
"total": 2,
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x4|vllm-mlu|text-generation": {
@@ -24077,9 +24121,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"参数/模板问题": 6
"参数/模板问题": 5
},
"failureCount": 6,
"failureCount": 5,
"failureRate": 1.0,
"framework": "unknown",
"lastPlatformFailureAt": null,
@@ -24091,8 +24135,8 @@
"successRate": 0.0,
"targetGpu": "hygon_k100-ai",
"taskType": "text-generation",
"total": 6,
"unresolvedFailureCount": 6
"total": 5,
"unresolvedFailureCount": 5
},
"hygon_k100-ai|vllm-patch-tokenizer|text-generation": {
"attributableFailureCount": 5,
@@ -24600,6 +24644,31 @@
"total": 1,
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x4|vllm-customized|text-generation|qwen3_5|none": {
"attributableFailureCount": 1,
"consecutiveFailures": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-customized",
"lastTerminalAt": "2026-09-29T17:54:53.270719+00:00",
"modelType": "qwen3_5",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x4",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Cambricon_mlu-370-x4|vllm-mlu|text-generation|llama|compressed-tensors": {
"attributableFailureCount": 0,
"consecutiveFailures": 0,
@@ -27582,14 +27651,14 @@
"unresolvedFailureCount": 1
},
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|llama|compressed-tensors|27": {
"attributableFailureCount": 1,
"attributableFailureCount": 2,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"decisionTotal": 2,
"failureBreakdown": {
"context_length": 1
"context_length": 2
},
"failureCount": 1,
"failureCount": 2,
"failureRate": 1.0,
"framework": "vllm_tokenizer_patch",
"loadSizeLog2Bucket": 27,
@@ -27602,7 +27671,7 @@
"successRate": 0.0,
"targetGpu": "Ascend_910-b3",
"taskType": "text-generation",
"total": 1,
"total": 2,
"unresolvedFailureCount": 0
},
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|llama|compressed-tensors|28": {
@@ -33828,6 +33897,30 @@
"total": 1,
"unresolvedFailureCount": 1
},
"Cambricon_mlu-370-x4|vllm-customized|text-generation|qwen3_5|none|33": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm-customized",
"loadSizeLog2Bucket": 33,
"modelType": "qwen3_5",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "none",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x4",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Cambricon_mlu-370-x4|vllm-customized|text-generation|stablelm|none|31": {
"attributableFailureCount": 0,
"decisionFailureRate": 0.0,
@@ -52868,20 +52961,20 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 17124,
"totalRecords": 17345,
"terminalRecords": 17126,
"totalRecords": 17347,
"totals": {
"attributableFailureCount": 6136,
"attributableFailureCount": 6138,
"decisionFailureRate": 0.8633,
"decisionSuccessRate": 0.1367,
"decisionTotal": 7108,
"decisionTotal": 7110,
"failureBreakdown": {
"ambiguous_runtime": 4451,
"architecture_compatibility": 212,
"attention_backend": 3,
"backend_operator": 128,
"context_length": 323,
"framework_architecture_unsupported": 2189,
"context_length": 324,
"framework_architecture_unsupported": 2190,
"memory_capacity": 1206,
"model_load": 550,
"platform_infrastructure": 942,
@@ -52892,14 +52985,14 @@
"日志缺失": 719,
"验证失败": 676
},
"failureCount": 16152,
"failureCount": 16154,
"failureRate": 0.9432,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 942,
"successCount": 972,
"successRate": 0.0568,
"total": 17124,
"total": 17126,
"unresolvedFailureCount": 9074
},
"warnings": [
@@ -52929,8 +53022,10 @@
"组合 Kunlunxin_p-800|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_mrv-100|unknown|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Kunlunxin_p-800|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x4|vllm-customized|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Mthreads_s4000|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -52953,11 +53048,10 @@
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x4|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Vastai_va16|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Kunlunxin_p-800|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
"组合 Vastai_va16|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 17345,
"summarizedRecords": 17347,
"version": 1
}

View File

@@ -14,13 +14,13 @@
100
],
"accounts": 12,
"activeScanned": 890,
"activeScanned": 888,
"ageCleanupMode": "admission_only",
"agePolicySkipped": {
"cleanupDisabled": true,
"reason": "admission_only"
},
"architectureBlockCount": 165,
"architectureBlockCount": 166,
"architectureFrameworkCatalog": {
"ascend_910-b3|text-generation": [
"llamacpp",
@@ -38,27 +38,11 @@
"vllm_fix_tokenizer",
"vllm_tokenizer_patch"
],
"cambricon_mlu-370-x4|text-generation": [
"vllm",
"vllm-customized",
"vllm-mlu"
],
"cambricon_mlu-370-x8|text-generation": [
"vllm",
"vllm-customized",
"vllm-mlu"
],
"hygon_k100-ai|text-generation": [
"llamacpp",
"vllm",
"vllm-patch-tokenizer"
],
"iluvatar_bi-100|text-generation": [
"transformers",
"vllm",
"vllm-patch-tokenizer",
"vllm_fix_tokenizer"
],
"iluvatar_bi-150|text-generation": [
"llamacpp",
"transformers",
@@ -67,19 +51,11 @@
"vllm_fix_tokenizer",
"vllm_tokenizer_patch"
],
"iluvatar_mrv-100|text-generation": [
"transformers",
"vllm"
],
"kunlunxin_p-800|text-generation": [
"vllm",
"vllm_fix_tokenizer",
"vllm_tokenizer_patch"
],
"metax_c-500|text-generation": [
"vllm",
"vllm_tokenizer_patch"
],
"mthreads_s4000|text-generation": [
"llamacpp",
"vllm",
@@ -89,10 +65,6 @@
"sglang",
"vllm",
"vllm_fix_tokenizer"
],
"vastai_va16|text-generation": [
"vllm",
"vllm_fix_tokenizer"
]
},
"architectureFrameworkCatalogErrors": {},
@@ -121,12 +93,12 @@
"z-lab/Qwen3.8-27B-DFlash2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/z-lab/Qwen3.8-27B-DFlash2-GGUF/resolve/master/config.json (status=404)"
},
"architectureModelConfigsComplete": 43,
"architectureOnly": false,
"architectureOnly": true,
"architecturePolicySkipped": {
"frameworkCatalogUnknown": 0,
"frameworkContextUnknown": 49,
"modelArchitectureUnknown": 146,
"noMatchingBlock": 727,
"noMatchingBlock": 725,
"partiallyBlockedFrameworkSet": 17,
"runningMatchedProtected": 0,
"submissionContextMismatch": 0,
@@ -166,15 +138,13 @@
"policyNoLongerAppliesTasks": [],
"recentModelDays": 7,
"recentModelReserveSlots": 5,
"repositorySizeErrors": {
"empero-ai/Qwen3.8-9B-GGUF": "recursive_repository_size_incomplete"
},
"repositorySizesComplete": 335,
"repositorySizeErrors": {},
"repositorySizesComplete": 0,
"skipped": {
"fitsKnownCapacity": 888,
"fitsKnownCapacity": 0,
"gpuCapacityUnknown": 0,
"repositorySizeUnknown": 2
"repositorySizeUnknown": 0
},
"stopErrors": [],
"uniqueModels": 336
"uniqueModels": 335
}

View File

@@ -34,6 +34,7 @@
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-22T10:15:00.769013+00:00", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-22T10:13:22+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4986624", "taskType": "text-generation", "verifyResult": null}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-22T10:15:00.769028+00:00", "modelId": "primitive-ai/Nex-N2.5-mini-NVFP4", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-22T10:13:22+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4986039", "taskType": "text-generation", "verifyResult": null}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-22T10:15:00.769035+00:00", "modelId": "mlx-community/Qwen3.5-122B-A10B-OptiQ-2bit-REAP-63B", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-22T10:13:22+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4979107", "taskType": "text-generation", "verifyResult": null}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm-customized", "lastSyncTime": "2026-09-29T17:54:53.270719+00:00", "modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "bb5d6992290b1b4fa8f4c37e884585a27bef87a1ebfb58e5ca8f098924550f05", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8674988799, "estimatedRequiredGiB": 9.718, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8695119291, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2415484144, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:speculative-decoding", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:qwen3-5", "custom_tag:9b", "custom_tag:6-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8695119291}, "outcome": "failed", "status": "success", "submitTime": "2026-09-22T05:27:11.098010+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "5012909", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["lfm2"], "framework": "vllm", "lastSyncTime": "2026-09-26T07:10:36.169057+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-22T05:27:11.087107+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "5012911", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_vl"], "framework": "vllm", "lastSyncTime": "2026-09-24T13:58:31.054930+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "8b3c0a82b858cf26b1a0f6075074234a66487d5734741dbbfd4743963b967faa", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-22T05:27:11.085801+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5012910", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5"], "framework": "vllm", "lastSyncTime": "2026-09-22T04:40:48.058549+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Quality", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-22T04:40:08+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4332831", "taskType": "text-generation", "verifyResult": -1}
@@ -297,4 +298,3 @@
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-21T02:55:04.446763+00:00", "modelId": "mlx-community/XYZ-Aquila-mini-OptiQ-4bit-REAP-18B", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T02:53:08+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4860051", "taskType": "text-generation", "verifyResult": null}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-21T02:55:04.446781+00:00", "modelId": "mlx-community/Ornith-1.5-35B-A3B-8bit", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T02:53:08+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4795634", "taskType": "text-generation", "verifyResult": null}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-21T02:55:04.446625+00:00", "modelId": "mlx-community/Agents-A1-OptiQ-4bit-REAP-19B", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T02:53:07+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4849006", "taskType": "text-generation", "verifyResult": null}
{"failReason": "参数/模板问题", "framework": "", "lastSyncTime": "2026-09-21T02:55:04.446646+00:00", "modelId": "mlx-community/KAT-Coder-V2.5-Dev-OptiQ-4bit", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-21T02:53:07+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4969661", "taskType": "text-generation", "verifyResult": null}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -2,24 +2,24 @@
"agentVersion": "2026.09.22.1",
"checksums": {
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
".modelhub_state/architecture_compatibility_blacklist.json": "3765c0b8cc37dbc4ab0853367dd78b5da141885fca1aa52fe4b39fd6f47b0dd4",
".modelhub_state/architecture_compatibility_blacklist.json": "69052afc91f2bb811a9a047a4405d79bad2fc314dcc810fc8120d78f2fb779ef",
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
".modelhub_state/market_intelligence.json": "2ec084fbae7d4a44d03b91fb6a30cbd0355ce44f72108d2bcc470e6e65bf54ed",
".modelhub_state/official_capabilities.json": "1acf5c13bdeb64858b3fc7a7c7cf4c9a5433afd9feef3e3bd3baa0546d2e6630",
".modelhub_state/outcome_checkpoint.json": "e9ea29ae6c08191e17ef80275e11113cd14342591f3972676f713ec472ad43a9",
".modelhub_state/queue_cleanup_latest.json": "49a85c11309f7edea696715997c353c42ca533519a99ab3a549526cccde3cad7",
".modelhub_state/recent_outcomes.jsonl": "81e91996ffb3668bc470e02ef1be371acb37d0b6aa5850bd175907421dd7d460",
".modelhub_state/recovery_active_tasks.jsonl": "baf4e60cd3f291bd320908dc1f58680e2118f753ba0ca85fdc8ab63a2ceff200",
".modelhub_state/recovery_intents.jsonl": "f295d8c004b998d4e0209fc857471aabe12849ba2044821017534bee41aa5bfc",
".modelhub_state/market_intelligence.json": "3ff3070c29e1526bafff1d85fe7dfeb4dbe40df7923662d3c838064fd2b662fe",
".modelhub_state/official_capabilities.json": "808396aac0a13eb952d0b4d623153a3ee9be73e9352635214c8ff548a326dc61",
".modelhub_state/outcome_checkpoint.json": "be57f950dcc2756929eec1446ef9f80930d465a811c3f9f7ba40dcda9cd7f4b4",
".modelhub_state/queue_cleanup_latest.json": "1b9a1363360f46e3333fceac79ff2c309befd678fffc2b42aff773c5823cf35e",
".modelhub_state/recent_outcomes.jsonl": "ea81f07d56492dc8073d6b8c3d83451fd1eb74f474036ebf59f262ff9711b54a",
".modelhub_state/recovery_active_tasks.jsonl": "31f91454ab698f5a0839e54ab2c37ad4dd4a17271d7ac23b6cffadd7393a99cf",
".modelhub_state/recovery_intents.jsonl": "a5cfd915d0bfbfa4b6f6791c760ecadb5ade9e2d878b3c17cc93ee5930ce03e6",
".modelhub_state/routing_intelligence.json": "895431dc6ea0fddce4c8a244da61cbaa7bdfb36cfdbd17e79f8e50e130dca87a",
".modelhub_state/submission_exclusions.jsonl": "1e62f6ba5bd2ba5f2f360bc134b44f7b69190230dd0e274919a87a73ea66761c",
".modelhub_state/worker_crashes.jsonl": "05bcde79076023fd20785b8270312f776430b95a2d4322ee53ef24e354c39f06",
"ledger/submissions.jsonl": "4e38a4e04b45021e3a767a110b76ae47d9635675eea4622c2e146306e2feb9a1",
"outcomes/submissions.jsonl": "1d385e9d0e0688d75933fd8c1df41a6174887206c5fcc25feb9642473b2f36e3"
"outcomes/submissions.jsonl": "3f794eddc38cdb71f8564a1232e8d2c402db9f8a51c0e863da47791a71d6811a"
},
"generation": 18476,
"generation": 18477,
"phase": "cycle",
"schemaVersion": 1,
"updatedAt": "2026-09-29T17:53:51.716839+00:00",
"updatedAt": "2026-09-29T17:58:00.051516+00:00",
"writerId": "9d958b95ce4146c9939da0f14b9201ff"
}

View File

@@ -299,7 +299,6 @@
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:30:54.600049+00:00", "modelId": "inceptionai/Jais-2-8B-Chat", "modelProfile": {"architectures": ["Jais2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16180861064, "estimatedRequiredGiB": 18.097, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "jais2", "modelscopeFileSize": 16193187697, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8090401280, "modelscopeTags": ["license:apache-2.0", "model_type:jais2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16193187697}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:32.237107+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969645", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T11:30:54.600410+00:00", "modelId": "BAAI/RoboBrain2.5-8B-NV", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918210, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918210}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:32.258565+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4969650", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:30:54.599773+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 54864980000, "estimatedRequiredGiB": 61.361, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 54904930780, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 54904930780}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:30:23.763321+00:00", "targetGpu": "MetaX_c-500", "taskId": "4969808", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:01:00.670531+00:00", "modelId": "neuralmagic/SmolLM-135M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "26bfee231695e882b96dee0191bc1dfec25bcad132af96a7ee5f0e26f3cb9fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 219850592, "estimatedRequiredGiB": 0.249, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 223236742, "modelscopeLicense": "apache-2.0", "modelscopeParams": 162826560, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 223236742}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:40:37.349301+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969958", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:01:00.670934+00:00", "modelId": "BAAI/RoboBrain2.5-8B-NV", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918210, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918210}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:48.995904+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970001", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.671009+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 54864980000, "estimatedRequiredGiB": 61.361, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 54904930780, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 54904930780}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:49.141735+00:00", "targetGpu": "Biren_166m", "taskId": "4970017", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.670603+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:49.147134+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4970014", "taskType": "text-generation", "verifyResult": null}
@@ -640,7 +639,6 @@
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T13:24:17.269090+00:00", "modelId": "nkkbr/Mini-K3-1H-decay-g1-v2_B", "modelProfile": {"architectures": [], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2030685704, "estimatedRequiredGiB": 2.279, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mini_k3", "modelscopeFileSize": null, "modelscopeLicense": null, "modelscopeParams": null, "modelscopeTags": ["model_type:mini_k3", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:kimi-k3", "custom_tag:pretraining", "custom_tag:mixture-of-experts", "custom_tag:linear-attention", "custom_tag:architecture-ablation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2039047610}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-22T05:22:37.692361+00:00", "targetGpu": "Biren_166m", "taskId": "5012842", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-22T13:24:17.269070+00:00", "modelId": "nkkbr/Mini-K3-1H-decay-g4-v2_B", "modelProfile": {"architectures": [], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2030769952, "estimatedRequiredGiB": 2.279, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mini_k3", "modelscopeFileSize": null, "modelscopeLicense": null, "modelscopeParams": null, "modelscopeTags": ["model_type:mini_k3", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:kimi-k3", "custom_tag:pretraining", "custom_tag:mixture-of-experts", "custom_tag:linear-attention", "custom_tag:architecture-ablation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2039132310}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-22T05:22:37.690620+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "5012841", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-22T13:24:17.269118+00:00", "modelId": "nkkbr/Mini-K3-1H-mamba2-n64-g4-v1_B", "modelProfile": {"architectures": [], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2009388256, "estimatedRequiredGiB": 2.25, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mini_k3", "modelscopeFileSize": null, "modelscopeLicense": null, "modelscopeParams": null, "modelscopeTags": ["model_type:mini_k3", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:kimi-k3", "custom_tag:mamba2", "custom_tag:pretraining", "custom_tag:mixture-of-experts", "custom_tag:architecture-ablation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2012995244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-22T05:22:37.689299+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "5012840", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-22T13:27:53.059266+00:00", "modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "bb5d6992290b1b4fa8f4c37e884585a27bef87a1ebfb58e5ca8f098924550f05", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8674988799, "estimatedRequiredGiB": 9.718, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8695119291, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2415484144, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:speculative-decoding", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:qwen3-5", "custom_tag:9b", "custom_tag:6-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8695119291}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-22T05:27:11.098010+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "5012909", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-22T13:30:17.459439+00:00", "modelId": "nkkbr/Mini-K3-1H-decay-g32-v2_C", "modelProfile": {"architectures": [], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2031556216, "estimatedRequiredGiB": 2.28, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mini_k3", "modelscopeFileSize": null, "modelscopeLicense": null, "modelscopeParams": null, "modelscopeTags": ["model_type:mini_k3", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:kimi-k3", "custom_tag:pretraining", "custom_tag:mixture-of-experts", "custom_tag:linear-attention", "custom_tag:architecture-ablation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2039918195}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-22T05:29:49.361193+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5012953", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-22T13:34:05.154403+00:00", "modelId": "Panyuqi/SpikingBrain-2.0-base-8k", "modelProfile": {"architectures": ["SSESWAMoBAForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 11110611058, "estimatedRequiredGiB": 12.435, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "sse_swa_moba", "modelscopeFileSize": 11126702266, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": null, "modelscopeTags": ["license:Apache License 2.0", "model_type:sse_swa_moba", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 11126702266}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-22T05:32:05.479742+00:00", "targetGpu": "MetaX_c-500", "taskId": "5012980", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-22T13:37:53.563477+00:00", "modelId": "nkkbr/Mini-K3-1H-kda-kernel-8-v2_D", "modelProfile": {"architectures": [], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2034915504, "estimatedRequiredGiB": 2.278, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mini_k3", "modelscopeFileSize": null, "modelscopeLicense": null, "modelscopeParams": null, "modelscopeTags": ["model_type:mini_k3", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:kimi-k3", "custom_tag:pretraining", "custom_tag:mixture-of-experts", "custom_tag:linear-attention", "custom_tag:architecture-ablation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2038530279}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-22T05:34:46.454174+00:00", "targetGpu": "Biren_166m", "taskId": "5013017", "taskType": "text-generation", "verifyResult": null}