state: generation 11676 (intent)
This commit is contained in:
@@ -1199,6 +1199,25 @@
|
|||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation"
|
"taskType": "text-generation"
|
||||||
},
|
},
|
||||||
|
"cambricon_mlu-370-x8|vllm|text-generation|model_type:qwen3_vl": {
|
||||||
|
"architectureSignature": "model_type:qwen3_vl",
|
||||||
|
"architectures": [],
|
||||||
|
"evidenceCount": 1,
|
||||||
|
"expiresAt": "2026-10-20T03:58:35.455197+00:00",
|
||||||
|
"framework": "vllm",
|
||||||
|
"latestFailureAt": "2026-09-20T03:58:35.455197+00:00",
|
||||||
|
"latestSuccessfulAt": null,
|
||||||
|
"matchType": "model_type",
|
||||||
|
"modelType": "qwen3_vl",
|
||||||
|
"sourceModelIds": [
|
||||||
|
"BAAI/RoboBrain2.5-8B-MT"
|
||||||
|
],
|
||||||
|
"sourceTaskIds": [
|
||||||
|
"4970358"
|
||||||
|
],
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
|
"taskType": "text-generation"
|
||||||
|
},
|
||||||
"cambricon_mlu-370-x8|vllm|text-generation|model_type:zaya": {
|
"cambricon_mlu-370-x8|vllm|text-generation|model_type:zaya": {
|
||||||
"architectureSignature": "model_type:zaya",
|
"architectureSignature": "model_type:zaya",
|
||||||
"architectures": [],
|
"architectures": [],
|
||||||
@@ -2768,9 +2787,9 @@
|
|||||||
"taskType": "text-generation"
|
"taskType": "text-generation"
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"generatedAt": "2026-09-21T19:54:03.658008+00:00",
|
"generatedAt": "2026-09-21T19:57:12.255712+00:00",
|
||||||
"summary": {
|
"summary": {
|
||||||
"activeBlockCount": 139,
|
"activeBlockCount": 140,
|
||||||
"byGpuFramework": {
|
"byGpuFramework": {
|
||||||
"Ascend_910-b3|vllm": 13,
|
"Ascend_910-b3|vllm": 13,
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
||||||
@@ -2778,7 +2797,7 @@
|
|||||||
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
||||||
"Biren_166m|vllm": 3,
|
"Biren_166m|vllm": 3,
|
||||||
"Cambricon_mlu-370-x4|vllm": 2,
|
"Cambricon_mlu-370-x4|vllm": 2,
|
||||||
"Cambricon_mlu-370-x8|vllm": 14,
|
"Cambricon_mlu-370-x8|vllm": 15,
|
||||||
"Cambricon_mlu-370-x8|vllm-customized": 2,
|
"Cambricon_mlu-370-x8|vllm-customized": 2,
|
||||||
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
||||||
"Iluvatar_bi-150|transformers": 4,
|
"Iluvatar_bi-150|transformers": 4,
|
||||||
|
|||||||
@@ -434,121 +434,121 @@
|
|||||||
}
|
}
|
||||||
},
|
},
|
||||||
"frameworkUpdatedAt": null,
|
"frameworkUpdatedAt": null,
|
||||||
"generatedAt": "2026-09-21T19:56:10.169860+00:00",
|
"generatedAt": "2026-09-21T19:57:30.800910+00:00",
|
||||||
"gpuStats": {
|
"gpuStats": {
|
||||||
"Ascend_910-b3": {
|
"Ascend_910-b3": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 929.2105263157895,
|
"backlogHours": 923.0980392156863,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Ascend_910-b3",
|
"gpu": "Ascend_910-b3",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 8,
|
"maxConcurrentTasks": 8,
|
||||||
"qualityFactor": 1.1979361109181033,
|
"qualityFactor": 1.5553566454179326,
|
||||||
"queueFactor": 0.9259023881211839,
|
"queueFactor": 0.9270529210260141,
|
||||||
"queueWeight": 0.9259023881211839,
|
"queueWeight": 0.9270529210260141,
|
||||||
"recentSuccess": 49,
|
"recentSuccess": 50,
|
||||||
"recentSuccessRate": 0.3223684210526316,
|
"recentSuccessRate": 0.32679738562091504,
|
||||||
"recentTerminal": 152,
|
"recentTerminal": 153,
|
||||||
"recentWilsonLowerBound": 0.2532349895684616,
|
"recentWilsonLowerBound": 0.25751021482630554,
|
||||||
"running": 8,
|
"running": 8,
|
||||||
"selectionWeight": 1.1091719059156753,
|
"selectionWeight": 1.441897921371917,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 25.333333333333332,
|
"throughputPerHour": 25.5,
|
||||||
"waiting": 23540
|
"waiting": 23539
|
||||||
},
|
},
|
||||||
"Ascend_910-b4": {
|
"Ascend_910-b4": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 1147.7432432432431,
|
"backlogHours": 1117.8157894736842,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Ascend_910-b4",
|
"gpu": "Ascend_910-b4",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 8,
|
"maxConcurrentTasks": 8,
|
||||||
"qualityFactor": 1.082389778474143,
|
"qualityFactor": 1.346809979880656,
|
||||||
"queueFactor": 0.8970271970857644,
|
"queueFactor": 0.9008161525985883,
|
||||||
"queueWeight": 0.8970271970857644,
|
"queueWeight": 0.9008161525985883,
|
||||||
"recentSuccess": 46,
|
"recentSuccess": 47,
|
||||||
"recentSuccessRate": 0.3108108108108108,
|
"recentSuccessRate": 0.3092105263157895,
|
||||||
"recentTerminal": 148,
|
"recentTerminal": 152,
|
||||||
"recentWilsonLowerBound": 0.24182488869231053,
|
"recentWilsonLowerBound": 0.24119850644698562,
|
||||||
"running": 8,
|
"running": 8,
|
||||||
"selectionWeight": 0.970933069138942,
|
"selectionWeight": 1.2132281843574746,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 24.666666666666668,
|
"throughputPerHour": 25.333333333333332,
|
||||||
"waiting": 28311
|
"waiting": 28318
|
||||||
},
|
},
|
||||||
"Biren_166m": {
|
"Biren_166m": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 165.68478260869566,
|
"backlogHours": 164.82162162162163,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Biren_166m",
|
"gpu": "Biren_166m",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 8,
|
"maxConcurrentTasks": 8,
|
||||||
"qualityFactor": 0.10780943019103657,
|
"qualityFactor": 0.11835048942061369,
|
||||||
"queueFactor": 1.1991953312870915,
|
"queueFactor": 1.2004375547181838,
|
||||||
"queueWeight": 1.1991953312870915,
|
"queueWeight": 1.2004375547181838,
|
||||||
"recentSuccess": 23,
|
"recentSuccess": 22,
|
||||||
"recentSuccessRate": 0.125,
|
"recentSuccessRate": 0.11891891891891893,
|
||||||
"recentTerminal": 184,
|
"recentTerminal": 185,
|
||||||
"recentWilsonLowerBound": 0.084756023564723,
|
"recentWilsonLowerBound": 0.07985694179220525,
|
||||||
"running": 8,
|
"running": 8,
|
||||||
"selectionWeight": 0.12928456535381266,
|
"selectionWeight": 0.14207237211978177,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 30.666666666666668,
|
"throughputPerHour": 30.833333333333332,
|
||||||
"waiting": 5081
|
"waiting": 5082
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x4": {
|
"Cambricon_mlu-370-x4": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 496.25274725274727,
|
"backlogHours": 496.31868131868134,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Cambricon_mlu-370-x4",
|
"gpu": "Cambricon_mlu-370-x4",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 7,
|
"maxConcurrentTasks": 7,
|
||||||
"qualityFactor": 0.05,
|
"qualityFactor": 0.05,
|
||||||
"queueFactor": 1.0172480916302264,
|
"queueFactor": 1.017484044376508,
|
||||||
"queueWeight": 1.0172480916302264,
|
"queueWeight": 1.017484044376508,
|
||||||
"recentSuccess": 14,
|
"recentSuccess": 14,
|
||||||
"recentSuccessRate": 0.07692307692307693,
|
"recentSuccessRate": 0.07692307692307693,
|
||||||
"recentTerminal": 182,
|
"recentTerminal": 182,
|
||||||
"recentWilsonLowerBound": 0.0463713966696411,
|
"recentWilsonLowerBound": 0.0463713966696411,
|
||||||
"running": 7,
|
"running": 7,
|
||||||
"selectionWeight": 0.050862404581511325,
|
"selectionWeight": 0.05087420221882541,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": false,
|
"submissionEligible": false,
|
||||||
"throughputPerHour": 30.333333333333332,
|
"throughputPerHour": 30.333333333333332,
|
||||||
"waiting": 15053
|
"waiting": 15055
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8": {
|
"Cambricon_mlu-370-x8": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 387.8450704225352,
|
"backlogHours": 369.94630872483225,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Cambricon_mlu-370-x8",
|
"gpu": "Cambricon_mlu-370-x8",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 7,
|
"maxConcurrentTasks": 7,
|
||||||
"qualityFactor": 0.21619746699449016,
|
"qualityFactor": 0.2706579774314779,
|
||||||
"queueFactor": 1.055561595771937,
|
"queueFactor": 1.0633369272014068,
|
||||||
"queueWeight": 1.055561595771937,
|
"queueWeight": 1.0633369272014068,
|
||||||
"recentSuccess": 24,
|
"recentSuccess": 25,
|
||||||
"recentSuccessRate": 0.16901408450704225,
|
"recentSuccessRate": 0.16778523489932887,
|
||||||
"recentTerminal": 142,
|
"recentTerminal": 149,
|
||||||
"recentWilsonLowerBound": 0.11628706631227818,
|
"recentWilsonLowerBound": 0.11630770351545354,
|
||||||
"running": 7,
|
"running": 7,
|
||||||
"selectionWeight": 0.22820974326255475,
|
"selectionWeight": 0.28780062204453544,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 23.666666666666668,
|
"throughputPerHour": 24.833333333333332,
|
||||||
"waiting": 9179
|
"waiting": 9187
|
||||||
},
|
},
|
||||||
"Iluvatar_bi-100": {
|
"Iluvatar_bi-100": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 38.66489361702128,
|
"backlogHours": 38.67914438502674,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Iluvatar_bi-100",
|
"gpu": "Iluvatar_bi-100",
|
||||||
@@ -558,129 +558,129 @@
|
|||||||
"queueFactor": 1.3,
|
"queueFactor": 1.3,
|
||||||
"queueWeight": 1.3,
|
"queueWeight": 1.3,
|
||||||
"recentSuccess": 21,
|
"recentSuccess": 21,
|
||||||
"recentSuccessRate": 0.05585106382978723,
|
"recentSuccessRate": 0.05614973262032086,
|
||||||
"recentTerminal": 376,
|
"recentTerminal": 374,
|
||||||
"recentWilsonLowerBound": 0.03681667599288761,
|
"recentWilsonLowerBound": 0.037015123146307144,
|
||||||
"running": 24,
|
"running": 26,
|
||||||
"selectionWeight": 0.065,
|
"selectionWeight": 0.065,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": false,
|
"submissionEligible": false,
|
||||||
"throughputPerHour": 62.666666666666664,
|
"throughputPerHour": 62.333333333333336,
|
||||||
"waiting": 2423
|
"waiting": 2411
|
||||||
},
|
},
|
||||||
"Iluvatar_bi-150": {
|
"Iluvatar_bi-150": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 111.1245283018868,
|
"backlogHours": 112.23664122137406,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Iluvatar_bi-150",
|
"gpu": "Iluvatar_bi-150",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 50,
|
"maxConcurrentTasks": 50,
|
||||||
"qualityFactor": 0.13312808001661144,
|
"qualityFactor": 0.1842073881825004,
|
||||||
"queueFactor": 1.273241638527949,
|
"queueFactor": 1.2716614373863284,
|
||||||
"queueWeight": 1.273241638527949,
|
"queueWeight": 1.2716614373863284,
|
||||||
"recentSuccess": 34,
|
"recentSuccess": 35,
|
||||||
"recentSuccessRate": 0.12830188679245283,
|
"recentSuccessRate": 0.13358778625954199,
|
||||||
"recentTerminal": 265,
|
"recentTerminal": 262,
|
||||||
"recentWilsonLowerBound": 0.0932852132110488,
|
"recentWilsonLowerBound": 0.09764447356554912,
|
||||||
"running": 7,
|
"running": 5,
|
||||||
"selectionWeight": 0.16950421473443025,
|
"selectionWeight": 0.23424943203333984,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 44.166666666666664,
|
"throughputPerHour": 43.666666666666664,
|
||||||
"waiting": 4908
|
"waiting": 4901
|
||||||
},
|
},
|
||||||
"Iluvatar_mrv-100": {
|
"Iluvatar_mrv-100": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 729.6818181818181,
|
"backlogHours": 740.9076923076923,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Iluvatar_mrv-100",
|
"gpu": "Iluvatar_mrv-100",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 2,
|
"maxConcurrentTasks": 2,
|
||||||
"qualityFactor": 0.05,
|
"qualityFactor": 0.05,
|
||||||
"queueFactor": 0.9600907680361976,
|
"queueFactor": 0.9581358392935574,
|
||||||
"queueWeight": 0.9600907680361976,
|
"queueWeight": 0.9581358392935574,
|
||||||
"recentSuccess": 12,
|
"recentSuccess": 11,
|
||||||
"recentSuccessRate": 0.09090909090909091,
|
"recentSuccessRate": 0.08461538461538462,
|
||||||
"recentTerminal": 132,
|
"recentTerminal": 130,
|
||||||
"recentWilsonLowerBound": 0.05276868798185061,
|
"recentWilsonLowerBound": 0.047903387035414836,
|
||||||
"running": 2,
|
"running": 3,
|
||||||
"selectionWeight": 0.04800453840180988,
|
"selectionWeight": 0.04790679196467787,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": false,
|
||||||
"throughputPerHour": 22.0,
|
"throughputPerHour": 21.666666666666668,
|
||||||
"waiting": 16053
|
"waiting": 16053
|
||||||
},
|
},
|
||||||
"Kunlunxin_p-800": {
|
"Kunlunxin_p-800": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 477.0676691729323,
|
"backlogHours": 469.8666666666667,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Kunlunxin_p-800",
|
"gpu": "Kunlunxin_p-800",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 8,
|
"maxConcurrentTasks": 8,
|
||||||
"qualityFactor": 0.27950214533336765,
|
"qualityFactor": 0.3032425396874696,
|
||||||
"queueFactor": 1.0232819760005352,
|
"queueFactor": 1.0258775016243378,
|
||||||
"queueWeight": 1.0232819760005352,
|
"queueWeight": 1.0258775016243378,
|
||||||
"recentSuccess": 25,
|
"recentSuccess": 24,
|
||||||
"recentSuccessRate": 0.18796992481203006,
|
"recentSuccessRate": 0.17777777777777778,
|
||||||
"recentTerminal": 133,
|
"recentTerminal": 135,
|
||||||
"recentWilsonLowerBound": 0.13068596058578308,
|
"recentWilsonLowerBound": 0.12247545696005419,
|
||||||
"running": 8,
|
"running": 8,
|
||||||
"selectionWeight": 0.2860095075731172,
|
"selectionWeight": 0.3110896990008004,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 22.166666666666668,
|
"throughputPerHour": 22.5,
|
||||||
"waiting": 10575
|
"waiting": 10572
|
||||||
},
|
},
|
||||||
"MetaX_c-500": {
|
"MetaX_c-500": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 608.9719626168225,
|
"backlogHours": 597.4678899082569,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "MetaX_c-500",
|
"gpu": "MetaX_c-500",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 4,
|
"maxConcurrentTasks": 4,
|
||||||
"qualityFactor": 0.508607241126908,
|
"qualityFactor": 0.7435278304250804,
|
||||||
"queueFactor": 0.9864900920111269,
|
"queueFactor": 0.9895654310354465,
|
||||||
"queueWeight": 0.9864900920111269,
|
"queueWeight": 0.9895654310354465,
|
||||||
"recentSuccess": 26,
|
"recentSuccess": 28,
|
||||||
"recentSuccessRate": 0.24299065420560748,
|
"recentSuccessRate": 0.25688073394495414,
|
||||||
"recentTerminal": 107,
|
"recentTerminal": 109,
|
||||||
"recentWilsonLowerBound": 0.17155744423820674,
|
"recentWilsonLowerBound": 0.18411863732206865,
|
||||||
"running": 4,
|
"running": 4,
|
||||||
"selectionWeight": 0.501736004096809,
|
"selectionWeight": 0.735769438001445,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 17.833333333333332,
|
"throughputPerHour": 18.166666666666668,
|
||||||
"waiting": 10860
|
"waiting": 10854
|
||||||
},
|
},
|
||||||
"Mthreads_s4000": {
|
"Mthreads_s4000": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 2574.3529411764707,
|
"backlogHours": 2680.2857142857147,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Mthreads_s4000",
|
"gpu": "Mthreads_s4000",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 4,
|
"maxConcurrentTasks": 4,
|
||||||
"qualityFactor": 2.041547163943165,
|
"qualityFactor": 2.4831426856003276,
|
||||||
"queueFactor": 0.7946613837998101,
|
"queueFactor": 0.7900681185897798,
|
||||||
"queueWeight": 0.7946613837998101,
|
"queueWeight": 0.7900681185897798,
|
||||||
"recentSuccess": 23,
|
"recentSuccess": 22,
|
||||||
"recentSuccessRate": 0.45098039215686275,
|
"recentSuccessRate": 0.4489795918367347,
|
||||||
"recentTerminal": 51,
|
"recentTerminal": 49,
|
||||||
"recentWilsonLowerBound": 0.3226730501113224,
|
"recentWilsonLowerBound": 0.31852624929636336,
|
||||||
"running": 4,
|
"running": 4,
|
||||||
"selectionWeight": 1.6223386943916531,
|
"selectionWeight": 1.9618518698022238,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 8.5,
|
"throughputPerHour": 8.166666666666666,
|
||||||
"waiting": 21882
|
"waiting": 21889
|
||||||
},
|
},
|
||||||
"Sunrise_pt-200-x1": {
|
"Sunrise_pt-200-x1": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 8617.09090909091,
|
"backlogHours": 7899.0,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Sunrise_pt-200-x1",
|
"gpu": "Sunrise_pt-200-x1",
|
||||||
@@ -690,64 +690,64 @@
|
|||||||
"queueFactor": 0.7,
|
"queueFactor": 0.7,
|
||||||
"queueWeight": 0.7,
|
"queueWeight": 0.7,
|
||||||
"recentSuccess": 7,
|
"recentSuccess": 7,
|
||||||
"recentSuccessRate": 0.6363636363636364,
|
"recentSuccessRate": 0.5833333333333334,
|
||||||
"recentTerminal": 11,
|
"recentTerminal": 12,
|
||||||
"recentWilsonLowerBound": 0.35379677689163724,
|
"recentWilsonLowerBound": 0.3195073356553728,
|
||||||
"running": 2,
|
"running": 2,
|
||||||
"selectionWeight": 1.75,
|
"selectionWeight": 1.75,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": false,
|
"submissionEligible": false,
|
||||||
"throughputPerHour": 1.8333333333333333,
|
"throughputPerHour": 2.0,
|
||||||
"waiting": 15798
|
"waiting": 15798
|
||||||
},
|
},
|
||||||
"Vastai_va16": {
|
"Vastai_va16": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 511.2352941176471,
|
"backlogHours": 520.5628742514971,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "Vastai_va16",
|
"gpu": "Vastai_va16",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 16,
|
"maxConcurrentTasks": 16,
|
||||||
"qualityFactor": 0.05,
|
"qualityFactor": 0.05,
|
||||||
"queueFactor": 1.0127195598592438,
|
"queueFactor": 1.010231072131324,
|
||||||
"queueWeight": 1.0127195598592438,
|
"queueWeight": 1.010231072131324,
|
||||||
"recentSuccess": 9,
|
"recentSuccess": 9,
|
||||||
"recentSuccessRate": 0.052941176470588235,
|
"recentSuccessRate": 0.05389221556886228,
|
||||||
"recentTerminal": 170,
|
"recentTerminal": 167,
|
||||||
"recentWilsonLowerBound": 0.028099063443071427,
|
"recentWilsonLowerBound": 0.02860843346782246,
|
||||||
"running": 16,
|
"running": 16,
|
||||||
"selectionWeight": 0.05063597799296219,
|
"selectionWeight": 0.0505115536065662,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": false,
|
"submissionEligible": false,
|
||||||
"throughputPerHour": 28.333333333333332,
|
"throughputPerHour": 27.833333333333332,
|
||||||
"waiting": 14485
|
"waiting": 14489
|
||||||
},
|
},
|
||||||
"hygon_k100-ai": {
|
"hygon_k100-ai": {
|
||||||
"available": true,
|
"available": true,
|
||||||
"backlogHours": 601.125,
|
"backlogHours": 593.6666666666666,
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
"error": null,
|
"error": null,
|
||||||
"gpu": "hygon_k100-ai",
|
"gpu": "hygon_k100-ai",
|
||||||
"healthFactor": 1.0,
|
"healthFactor": 1.0,
|
||||||
"maxConcurrentTasks": 6,
|
"maxConcurrentTasks": 6,
|
||||||
"qualityFactor": 0.0761241074866858,
|
"qualityFactor": 0.10731637073610716,
|
||||||
"queueFactor": 0.9884110770796999,
|
"queueFactor": 0.9905132768694215,
|
||||||
"queueWeight": 0.9884110770796999,
|
"queueWeight": 0.9905132768694215,
|
||||||
"recentSuccess": 18,
|
"recentSuccess": 19,
|
||||||
"recentSuccessRate": 0.1125,
|
"recentSuccessRate": 0.11728395061728394,
|
||||||
"recentTerminal": 160,
|
"recentTerminal": 162,
|
||||||
"recentWilsonLowerBound": 0.07235575112573997,
|
"recentWilsonLowerBound": 0.07638228346042165,
|
||||||
"running": 6,
|
"running": 6,
|
||||||
"selectionWeight": 0.07524191107264595,
|
"selectionWeight": 0.10629829003955518,
|
||||||
"stale": false,
|
"stale": false,
|
||||||
"submissionEligible": true,
|
"submissionEligible": true,
|
||||||
"throughputPerHour": 26.666666666666668,
|
"throughputPerHour": 27.0,
|
||||||
"waiting": 16030
|
"waiting": 16029
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"queueAttemptedAt": "2026-09-21T19:47:15.264791+00:00",
|
"queueAttemptedAt": "2026-09-21T19:57:30.800910+00:00",
|
||||||
"queueError": null,
|
"queueError": null,
|
||||||
"queueUpdatedAt": "2026-09-21T19:47:15.264791+00:00",
|
"queueUpdatedAt": "2026-09-21T19:57:30.800910+00:00",
|
||||||
"supportedGpus": [
|
"supportedGpus": [
|
||||||
"Vastai_va16",
|
"Vastai_va16",
|
||||||
"Kunlunxin_p-800",
|
"Kunlunxin_p-800",
|
||||||
|
|||||||
@@ -1,5 +1,5 @@
|
|||||||
{
|
{
|
||||||
"catalogUpdatedAt": "2026-09-21T19:56:10.169860+00:00",
|
"catalogUpdatedAt": "2026-09-21T19:57:30.800910+00:00",
|
||||||
"configuredTaskTypes": [
|
"configuredTaskTypes": [
|
||||||
"text-generation"
|
"text-generation"
|
||||||
],
|
],
|
||||||
@@ -56,7 +56,7 @@
|
|||||||
"time-series-forecasting"
|
"time-series-forecasting"
|
||||||
],
|
],
|
||||||
"errors": [],
|
"errors": [],
|
||||||
"generatedAt": "2026-09-21T19:56:10.169860+00:00",
|
"generatedAt": "2026-09-21T19:57:30.800910+00:00",
|
||||||
"gpuCatalog": {
|
"gpuCatalog": {
|
||||||
"Ascend_910-b3": {
|
"Ascend_910-b3": {
|
||||||
"canVerify": true,
|
"canVerify": true,
|
||||||
@@ -6886,6 +6886,6 @@
|
|||||||
"updateTime": "2025-12-22 08:59:53"
|
"updateTime": "2025-12-22 08:59:53"
|
||||||
}
|
}
|
||||||
],
|
],
|
||||||
"taskTreeUpdatedAt": "2026-09-21T19:56:10.169860+00:00",
|
"taskTreeUpdatedAt": "2026-09-21T19:57:30.800910+00:00",
|
||||||
"version": 1
|
"version": 1
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -1,6 +1,6 @@
|
|||||||
{
|
{
|
||||||
"generatedAt": "2026-09-21T19:52:34.453108+00:00",
|
"generatedAt": "2026-09-21T19:57:12.166601+00:00",
|
||||||
"lastSyncTime": "2026-09-21T19:52:34.157602+00:00",
|
"lastSyncTime": "2026-09-21T19:57:11.769139+00:00",
|
||||||
"recentLimit": 300,
|
"recentLimit": 300,
|
||||||
"report": {
|
"report": {
|
||||||
"architectureCompatibilityBlocks": {
|
"architectureCompatibilityBlocks": {
|
||||||
@@ -1203,6 +1203,25 @@
|
|||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation"
|
"taskType": "text-generation"
|
||||||
},
|
},
|
||||||
|
"cambricon_mlu-370-x8|vllm|text-generation|model_type:qwen3_vl": {
|
||||||
|
"architectureSignature": "model_type:qwen3_vl",
|
||||||
|
"architectures": [],
|
||||||
|
"evidenceCount": 1,
|
||||||
|
"expiresAt": "2026-10-20T03:58:35.455197+00:00",
|
||||||
|
"framework": "vllm",
|
||||||
|
"latestFailureAt": "2026-09-20T03:58:35.455197+00:00",
|
||||||
|
"latestSuccessfulAt": null,
|
||||||
|
"matchType": "model_type",
|
||||||
|
"modelType": "qwen3_vl",
|
||||||
|
"sourceModelIds": [
|
||||||
|
"BAAI/RoboBrain2.5-8B-MT"
|
||||||
|
],
|
||||||
|
"sourceTaskIds": [
|
||||||
|
"4970358"
|
||||||
|
],
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
|
"taskType": "text-generation"
|
||||||
|
},
|
||||||
"cambricon_mlu-370-x8|vllm|text-generation|model_type:zaya": {
|
"cambricon_mlu-370-x8|vllm|text-generation|model_type:zaya": {
|
||||||
"architectureSignature": "model_type:zaya",
|
"architectureSignature": "model_type:zaya",
|
||||||
"architectures": [],
|
"architectures": [],
|
||||||
@@ -2773,7 +2792,7 @@
|
|||||||
}
|
}
|
||||||
},
|
},
|
||||||
"architectureCompatibilitySummary": {
|
"architectureCompatibilitySummary": {
|
||||||
"activeBlockCount": 139,
|
"activeBlockCount": 140,
|
||||||
"byGpuFramework": {
|
"byGpuFramework": {
|
||||||
"Ascend_910-b3|vllm": 13,
|
"Ascend_910-b3|vllm": 13,
|
||||||
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
||||||
@@ -2781,7 +2800,7 @@
|
|||||||
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
"Ascend_910-b4|vllm_tokenizer_patch": 5,
|
||||||
"Biren_166m|vllm": 3,
|
"Biren_166m|vllm": 3,
|
||||||
"Cambricon_mlu-370-x4|vllm": 2,
|
"Cambricon_mlu-370-x4|vllm": 2,
|
||||||
"Cambricon_mlu-370-x8|vllm": 14,
|
"Cambricon_mlu-370-x8|vllm": 15,
|
||||||
"Cambricon_mlu-370-x8|vllm-customized": 2,
|
"Cambricon_mlu-370-x8|vllm-customized": 2,
|
||||||
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
"Cambricon_mlu-370-x8|vllm-mlu": 5,
|
||||||
"Iluvatar_bi-150|transformers": 4,
|
"Iluvatar_bi-150|transformers": 4,
|
||||||
@@ -3077,13 +3096,13 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 24,
|
"decisionTotal": 24,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 45,
|
"ambiguous_runtime": 46,
|
||||||
"framework_architecture_unsupported": 23,
|
"framework_architecture_unsupported": 23,
|
||||||
"platform_infrastructure": 1,
|
"platform_infrastructure": 1,
|
||||||
"repository_structure": 1,
|
"repository_structure": 1,
|
||||||
"参数/模板问题": 6
|
"参数/模板问题": 6
|
||||||
},
|
},
|
||||||
"failureCount": 76,
|
"failureCount": 77,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm_tokenizer_patch",
|
"framework": "vllm_tokenizer_patch",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
@@ -3093,8 +3112,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Ascend_910-b4",
|
"targetGpu": "Ascend_910-b4",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 76,
|
"total": 77,
|
||||||
"unresolvedFailureCount": 51
|
"unresolvedFailureCount": 52
|
||||||
},
|
},
|
||||||
"Ascend_910-b4|vllm|text-generation": {
|
"Ascend_910-b4|vllm|text-generation": {
|
||||||
"attributableFailureCount": 204,
|
"attributableFailureCount": 204,
|
||||||
@@ -3446,37 +3465,37 @@
|
|||||||
"decisionSuccessRate": 0.48,
|
"decisionSuccessRate": 0.48,
|
||||||
"decisionTotal": 25,
|
"decisionTotal": 25,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 39,
|
"ambiguous_runtime": 40,
|
||||||
"framework_architecture_unsupported": 11,
|
"framework_architecture_unsupported": 11,
|
||||||
"model_load": 1,
|
"model_load": 1,
|
||||||
"tokenizer_compatibility": 1,
|
"tokenizer_compatibility": 1,
|
||||||
"参数/模板问题": 3
|
"参数/模板问题": 3
|
||||||
},
|
},
|
||||||
"failureCount": 55,
|
"failureCount": 56,
|
||||||
"failureRate": 0.8209,
|
"failureRate": 0.8235,
|
||||||
"framework": "vllm-mlu",
|
"framework": "vllm-mlu",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 0,
|
"platformFailureCount": 0,
|
||||||
"successCount": 12,
|
"successCount": 12,
|
||||||
"successRate": 0.1791,
|
"successRate": 0.1765,
|
||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 67,
|
"total": 68,
|
||||||
"unresolvedFailureCount": 42
|
"unresolvedFailureCount": 43
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm|text-generation": {
|
"Cambricon_mlu-370-x8|vllm|text-generation": {
|
||||||
"attributableFailureCount": 19,
|
"attributableFailureCount": 20,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 19,
|
"decisionTotal": 20,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 19,
|
"ambiguous_runtime": 19,
|
||||||
"framework_architecture_unsupported": 18,
|
"framework_architecture_unsupported": 19,
|
||||||
"memory_capacity": 1,
|
"memory_capacity": 1,
|
||||||
"参数/模板问题": 12
|
"参数/模板问题": 12
|
||||||
},
|
},
|
||||||
"failureCount": 50,
|
"failureCount": 51,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
@@ -3486,7 +3505,7 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 50,
|
"total": 51,
|
||||||
"unresolvedFailureCount": 31
|
"unresolvedFailureCount": 31
|
||||||
},
|
},
|
||||||
"Iluvatar_bi-100|transformers|text-generation": {
|
"Iluvatar_bi-100|transformers|text-generation": {
|
||||||
@@ -5421,17 +5440,17 @@
|
|||||||
"unresolvedFailureCount": 6446
|
"unresolvedFailureCount": 6446
|
||||||
},
|
},
|
||||||
"vllm": {
|
"vllm": {
|
||||||
"attributableFailureCount": 3542,
|
"attributableFailureCount": 3543,
|
||||||
"decisionFailureRate": 0.9758,
|
"decisionFailureRate": 0.9758,
|
||||||
"decisionSuccessRate": 0.0242,
|
"decisionSuccessRate": 0.0242,
|
||||||
"decisionTotal": 3630,
|
"decisionTotal": 3631,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 1497,
|
"ambiguous_runtime": 1497,
|
||||||
"architecture_compatibility": 112,
|
"architecture_compatibility": 112,
|
||||||
"attention_backend": 1,
|
"attention_backend": 1,
|
||||||
"backend_operator": 84,
|
"backend_operator": 84,
|
||||||
"context_length": 161,
|
"context_length": 161,
|
||||||
"framework_architecture_unsupported": 1309,
|
"framework_architecture_unsupported": 1310,
|
||||||
"memory_capacity": 722,
|
"memory_capacity": 722,
|
||||||
"model_load": 207,
|
"model_load": 207,
|
||||||
"platform_infrastructure": 865,
|
"platform_infrastructure": 865,
|
||||||
@@ -5440,14 +5459,14 @@
|
|||||||
"tokenizer_compatibility": 412,
|
"tokenizer_compatibility": 412,
|
||||||
"参数/模板问题": 46
|
"参数/模板问题": 46
|
||||||
},
|
},
|
||||||
"failureCount": 5950,
|
"failureCount": 5951,
|
||||||
"failureRate": 0.9854,
|
"failureRate": 0.9854,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 865,
|
"platformFailureCount": 865,
|
||||||
"successCount": 88,
|
"successCount": 88,
|
||||||
"successRate": 0.0146,
|
"successRate": 0.0146,
|
||||||
"total": 6038,
|
"total": 6039,
|
||||||
"unresolvedFailureCount": 1543
|
"unresolvedFailureCount": 1543
|
||||||
},
|
},
|
||||||
"vllm-customized": {
|
"vllm-customized": {
|
||||||
@@ -5477,22 +5496,22 @@
|
|||||||
"decisionSuccessRate": 0.4878,
|
"decisionSuccessRate": 0.4878,
|
||||||
"decisionTotal": 41,
|
"decisionTotal": 41,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 58,
|
"ambiguous_runtime": 59,
|
||||||
"framework_architecture_unsupported": 18,
|
"framework_architecture_unsupported": 18,
|
||||||
"model_load": 2,
|
"model_load": 2,
|
||||||
"tokenizer_compatibility": 1,
|
"tokenizer_compatibility": 1,
|
||||||
"参数/模板问题": 3,
|
"参数/模板问题": 3,
|
||||||
"验证失败": 1
|
"验证失败": 1
|
||||||
},
|
},
|
||||||
"failureCount": 83,
|
"failureCount": 84,
|
||||||
"failureRate": 0.8058,
|
"failureRate": 0.8077,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 0,
|
"platformFailureCount": 0,
|
||||||
"successCount": 20,
|
"successCount": 20,
|
||||||
"successRate": 0.1942,
|
"successRate": 0.1923,
|
||||||
"total": 103,
|
"total": 104,
|
||||||
"unresolvedFailureCount": 62
|
"unresolvedFailureCount": 63
|
||||||
},
|
},
|
||||||
"vllm-patch-tokenizer": {
|
"vllm-patch-tokenizer": {
|
||||||
"attributableFailureCount": 31,
|
"attributableFailureCount": 31,
|
||||||
@@ -5572,7 +5591,7 @@
|
|||||||
"decisionSuccessRate": 0.0263,
|
"decisionSuccessRate": 0.0263,
|
||||||
"decisionTotal": 76,
|
"decisionTotal": 76,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 80,
|
"ambiguous_runtime": 81,
|
||||||
"backend_operator": 1,
|
"backend_operator": 1,
|
||||||
"context_length": 3,
|
"context_length": 3,
|
||||||
"framework_architecture_unsupported": 63,
|
"framework_architecture_unsupported": 63,
|
||||||
@@ -5583,18 +5602,18 @@
|
|||||||
"tokenizer_compatibility": 1,
|
"tokenizer_compatibility": 1,
|
||||||
"参数/模板问题": 27
|
"参数/模板问题": 27
|
||||||
},
|
},
|
||||||
"failureCount": 182,
|
"failureCount": 183,
|
||||||
"failureRate": 0.9891,
|
"failureRate": 0.9892,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 1,
|
"platformFailureCount": 1,
|
||||||
"successCount": 2,
|
"successCount": 2,
|
||||||
"successRate": 0.0109,
|
"successRate": 0.0108,
|
||||||
"total": 184,
|
"total": 185,
|
||||||
"unresolvedFailureCount": 107
|
"unresolvedFailureCount": 108
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"generatedAt": "2026-09-21T19:52:34.442165+00:00",
|
"generatedAt": "2026-09-21T19:57:12.116575+00:00",
|
||||||
"gpuSummaries": {
|
"gpuSummaries": {
|
||||||
"Ascend_910-b3": {
|
"Ascend_910-b3": {
|
||||||
"attributableFailureCount": 99,
|
"attributableFailureCount": 99,
|
||||||
@@ -5628,7 +5647,7 @@
|
|||||||
"decisionSuccessRate": 0.1636,
|
"decisionSuccessRate": 0.1636,
|
||||||
"decisionTotal": 385,
|
"decisionTotal": 385,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 181,
|
"ambiguous_runtime": 182,
|
||||||
"context_length": 1,
|
"context_length": 1,
|
||||||
"framework_architecture_unsupported": 178,
|
"framework_architecture_unsupported": 178,
|
||||||
"memory_capacity": 19,
|
"memory_capacity": 19,
|
||||||
@@ -5640,15 +5659,15 @@
|
|||||||
"日志缺失": 14,
|
"日志缺失": 14,
|
||||||
"验证失败": 175
|
"验证失败": 175
|
||||||
},
|
},
|
||||||
"failureCount": 962,
|
"failureCount": 963,
|
||||||
"failureRate": 0.9385,
|
"failureRate": 0.9386,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 2,
|
"platformFailureCount": 2,
|
||||||
"successCount": 63,
|
"successCount": 63,
|
||||||
"successRate": 0.0615,
|
"successRate": 0.0614,
|
||||||
"total": 1025,
|
"total": 1026,
|
||||||
"unresolvedFailureCount": 638
|
"unresolvedFailureCount": 639
|
||||||
},
|
},
|
||||||
"Biren_166m": {
|
"Biren_166m": {
|
||||||
"attributableFailureCount": 186,
|
"attributableFailureCount": 186,
|
||||||
@@ -5710,28 +5729,28 @@
|
|||||||
"unresolvedFailureCount": 774
|
"unresolvedFailureCount": 774
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8": {
|
"Cambricon_mlu-370-x8": {
|
||||||
"attributableFailureCount": 50,
|
"attributableFailureCount": 51,
|
||||||
"decisionFailureRate": 0.7812,
|
"decisionFailureRate": 0.7846,
|
||||||
"decisionSuccessRate": 0.2188,
|
"decisionSuccessRate": 0.2154,
|
||||||
"decisionTotal": 64,
|
"decisionTotal": 65,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 108,
|
"ambiguous_runtime": 109,
|
||||||
"framework_architecture_unsupported": 41,
|
"framework_architecture_unsupported": 42,
|
||||||
"memory_capacity": 4,
|
"memory_capacity": 4,
|
||||||
"model_load": 3,
|
"model_load": 3,
|
||||||
"tokenizer_compatibility": 2,
|
"tokenizer_compatibility": 2,
|
||||||
"参数/模板问题": 34,
|
"参数/模板问题": 34,
|
||||||
"验证失败": 22
|
"验证失败": 22
|
||||||
},
|
},
|
||||||
"failureCount": 214,
|
"failureCount": 216,
|
||||||
"failureRate": 0.9386,
|
"failureRate": 0.9391,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 0,
|
"platformFailureCount": 0,
|
||||||
"successCount": 14,
|
"successCount": 14,
|
||||||
"successRate": 0.0614,
|
"successRate": 0.0609,
|
||||||
"total": 228,
|
"total": 230,
|
||||||
"unresolvedFailureCount": 164
|
"unresolvedFailureCount": 165
|
||||||
},
|
},
|
||||||
"Iluvatar_bi-100": {
|
"Iluvatar_bi-100": {
|
||||||
"attributableFailureCount": 194,
|
"attributableFailureCount": 194,
|
||||||
@@ -7460,10 +7479,10 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 2,
|
"ambiguous_runtime": 3,
|
||||||
"platform_infrastructure": 1
|
"platform_infrastructure": 1
|
||||||
},
|
},
|
||||||
"failureCount": 3,
|
"failureCount": 4,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm_tokenizer_patch",
|
"framework": "vllm_tokenizer_patch",
|
||||||
"modelType": "qwen2",
|
"modelType": "qwen2",
|
||||||
@@ -7475,8 +7494,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Ascend_910-b4",
|
"targetGpu": "Ascend_910-b4",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 3,
|
"total": 4,
|
||||||
"unresolvedFailureCount": 2
|
"unresolvedFailureCount": 3
|
||||||
},
|
},
|
||||||
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|qwen3_5_moe|compressed-tensors": {
|
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|qwen3_5_moe|compressed-tensors": {
|
||||||
"attributableFailureCount": 3,
|
"attributableFailureCount": 3,
|
||||||
@@ -9877,9 +9896,9 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 1
|
"ambiguous_runtime": 2
|
||||||
},
|
},
|
||||||
"failureCount": 1,
|
"failureCount": 2,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm-mlu",
|
"framework": "vllm-mlu",
|
||||||
"modelType": "lfm2",
|
"modelType": "lfm2",
|
||||||
@@ -9891,8 +9910,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 1,
|
"total": 2,
|
||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 2
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|llama|awq": {
|
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|llama|awq": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -10884,6 +10903,29 @@
|
|||||||
"total": 2,
|
"total": 2,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3_vl|none": {
|
||||||
|
"attributableFailureCount": 1,
|
||||||
|
"decisionFailureRate": 1.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 1,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"framework_architecture_unsupported": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"modelType": "qwen3_vl",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 0
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3|none": {
|
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3|none": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
@@ -17463,12 +17505,12 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 6,
|
"decisionTotal": 6,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 8,
|
"ambiguous_runtime": 7,
|
||||||
"framework_architecture_unsupported": 1,
|
"framework_architecture_unsupported": 1,
|
||||||
"model_load": 4,
|
"model_load": 4,
|
||||||
"tokenizer_compatibility": 1
|
"tokenizer_compatibility": 1
|
||||||
},
|
},
|
||||||
"failureCount": 14,
|
"failureCount": 13,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"lastPlatformFailureAt": null,
|
"lastPlatformFailureAt": null,
|
||||||
@@ -17480,8 +17522,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Biren_166m",
|
"targetGpu": "Biren_166m",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 14,
|
"total": 13,
|
||||||
"unresolvedFailureCount": 8
|
"unresolvedFailureCount": 7
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x4|unknown|text-generation": {
|
"Cambricon_mlu-370-x4|unknown|text-generation": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -17569,9 +17611,9 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 6
|
"ambiguous_runtime": 7
|
||||||
},
|
},
|
||||||
"failureCount": 6,
|
"failureCount": 7,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm-mlu",
|
"framework": "vllm-mlu",
|
||||||
"lastPlatformFailureAt": null,
|
"lastPlatformFailureAt": null,
|
||||||
@@ -17583,8 +17625,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 6,
|
"total": 7,
|
||||||
"unresolvedFailureCount": 6
|
"unresolvedFailureCount": 7
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm|text-generation": {
|
"Cambricon_mlu-370-x8|vllm|text-generation": {
|
||||||
"attributableFailureCount": 6,
|
"attributableFailureCount": 6,
|
||||||
@@ -18640,9 +18682,9 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 2
|
"ambiguous_runtime": 1
|
||||||
},
|
},
|
||||||
"failureCount": 2,
|
"failureCount": 1,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm",
|
"framework": "vllm",
|
||||||
"lastTerminalAt": "2026-09-21T05:07:16.468545+00:00",
|
"lastTerminalAt": "2026-09-21T05:07:16.468545+00:00",
|
||||||
@@ -18655,8 +18697,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Biren_166m",
|
"targetGpu": "Biren_166m",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 2,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 2
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
"Biren_166m|vllm|text-generation|lfm2|none": {
|
"Biren_166m|vllm|text-generation|lfm2|none": {
|
||||||
"attributableFailureCount": 2,
|
"attributableFailureCount": 2,
|
||||||
@@ -19019,12 +19061,12 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 1
|
"ambiguous_runtime": 2
|
||||||
},
|
},
|
||||||
"failureCount": 1,
|
"failureCount": 2,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm-mlu",
|
"framework": "vllm-mlu",
|
||||||
"lastTerminalAt": "2026-09-21T08:44:09.587559+00:00",
|
"lastTerminalAt": "2026-09-21T19:57:11.769116+00:00",
|
||||||
"modelType": "lfm2",
|
"modelType": "lfm2",
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
@@ -19034,8 +19076,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 1,
|
"total": 2,
|
||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 2
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|phi3|compressed-tensors": {
|
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|phi3|compressed-tensors": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -19487,6 +19529,31 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 1
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3_vl|none": {
|
||||||
|
"attributableFailureCount": 1,
|
||||||
|
"consecutiveFailures": 1,
|
||||||
|
"decisionFailureRate": 1.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 1,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"framework_architecture_unsupported": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"lastTerminalAt": "2026-09-21T19:57:11.769130+00:00",
|
||||||
|
"modelType": "qwen3_vl",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 0
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3|none": {
|
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3|none": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"consecutiveFailures": 1,
|
"consecutiveFailures": 1,
|
||||||
@@ -19962,31 +20029,6 @@
|
|||||||
"total": 1,
|
"total": 1,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
"Iluvatar_bi-150|transformers|text-generation|mpt|none": {
|
|
||||||
"attributableFailureCount": 0,
|
|
||||||
"consecutiveFailures": 0,
|
|
||||||
"decisionFailureRate": 0.0,
|
|
||||||
"decisionSuccessRate": 0.0,
|
|
||||||
"decisionTotal": 0,
|
|
||||||
"failureBreakdown": {
|
|
||||||
"ambiguous_runtime": 1
|
|
||||||
},
|
|
||||||
"failureCount": 1,
|
|
||||||
"failureRate": 1.0,
|
|
||||||
"framework": "transformers",
|
|
||||||
"lastTerminalAt": "2026-09-20T14:14:39.678925+00:00",
|
|
||||||
"modelType": "mpt",
|
|
||||||
"pendingCount": 0,
|
|
||||||
"pendingRate": 0.0,
|
|
||||||
"platformFailureCount": 0,
|
|
||||||
"quantizationMethod": "none",
|
|
||||||
"successCount": 0,
|
|
||||||
"successRate": 0.0,
|
|
||||||
"targetGpu": "Iluvatar_bi-150",
|
|
||||||
"taskType": "text-generation",
|
|
||||||
"total": 1,
|
|
||||||
"unresolvedFailureCount": 1
|
|
||||||
},
|
|
||||||
"Iluvatar_bi-150|transformers|text-generation|nemotron_h|compressed-tensors": {
|
"Iluvatar_bi-150|transformers|text-generation|nemotron_h|compressed-tensors": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"consecutiveFailures": 1,
|
"consecutiveFailures": 1,
|
||||||
@@ -23497,9 +23539,9 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 1
|
"ambiguous_runtime": 2
|
||||||
},
|
},
|
||||||
"failureCount": 1,
|
"failureCount": 2,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm_tokenizer_patch",
|
"framework": "vllm_tokenizer_patch",
|
||||||
"loadSizeLog2Bucket": 29,
|
"loadSizeLog2Bucket": 29,
|
||||||
@@ -23512,8 +23554,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Ascend_910-b4",
|
"targetGpu": "Ascend_910-b4",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 1,
|
"total": 2,
|
||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 2
|
||||||
},
|
},
|
||||||
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|qwen2|compressed-tensors|31": {
|
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|qwen2|compressed-tensors|31": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -26943,9 +26985,9 @@
|
|||||||
"decisionSuccessRate": 0.0,
|
"decisionSuccessRate": 0.0,
|
||||||
"decisionTotal": 0,
|
"decisionTotal": 0,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 1
|
"ambiguous_runtime": 2
|
||||||
},
|
},
|
||||||
"failureCount": 1,
|
"failureCount": 2,
|
||||||
"failureRate": 1.0,
|
"failureRate": 1.0,
|
||||||
"framework": "vllm-mlu",
|
"framework": "vllm-mlu",
|
||||||
"loadSizeLog2Bucket": 29,
|
"loadSizeLog2Bucket": 29,
|
||||||
@@ -26958,8 +27000,8 @@
|
|||||||
"successRate": 0.0,
|
"successRate": 0.0,
|
||||||
"targetGpu": "Cambricon_mlu-370-x8",
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
"taskType": "text-generation",
|
"taskType": "text-generation",
|
||||||
"total": 1,
|
"total": 2,
|
||||||
"unresolvedFailureCount": 1
|
"unresolvedFailureCount": 2
|
||||||
},
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|llama|awq|30": {
|
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|llama|awq|30": {
|
||||||
"attributableFailureCount": 0,
|
"attributableFailureCount": 0,
|
||||||
@@ -28373,6 +28415,30 @@
|
|||||||
"total": 2,
|
"total": 2,
|
||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
},
|
},
|
||||||
|
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3_vl|none|34": {
|
||||||
|
"attributableFailureCount": 1,
|
||||||
|
"decisionFailureRate": 1.0,
|
||||||
|
"decisionSuccessRate": 0.0,
|
||||||
|
"decisionTotal": 1,
|
||||||
|
"failureBreakdown": {
|
||||||
|
"framework_architecture_unsupported": 1
|
||||||
|
},
|
||||||
|
"failureCount": 1,
|
||||||
|
"failureRate": 1.0,
|
||||||
|
"framework": "vllm",
|
||||||
|
"loadSizeLog2Bucket": 34,
|
||||||
|
"modelType": "qwen3_vl",
|
||||||
|
"pendingCount": 0,
|
||||||
|
"pendingRate": 0.0,
|
||||||
|
"platformFailureCount": 0,
|
||||||
|
"quantizationMethod": "none",
|
||||||
|
"successCount": 0,
|
||||||
|
"successRate": 0.0,
|
||||||
|
"targetGpu": "Cambricon_mlu-370-x8",
|
||||||
|
"taskType": "text-generation",
|
||||||
|
"total": 1,
|
||||||
|
"unresolvedFailureCount": 0
|
||||||
|
},
|
||||||
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3|none|29": {
|
"Cambricon_mlu-370-x8|vllm|text-generation|qwen3|none|29": {
|
||||||
"attributableFailureCount": 1,
|
"attributableFailureCount": 1,
|
||||||
"decisionFailureRate": 1.0,
|
"decisionFailureRate": 1.0,
|
||||||
@@ -37637,20 +37703,20 @@
|
|||||||
"unresolvedFailureCount": 0
|
"unresolvedFailureCount": 0
|
||||||
}
|
}
|
||||||
},
|
},
|
||||||
"terminalRecords": 16388,
|
"terminalRecords": 16391,
|
||||||
"totalRecords": 16582,
|
"totalRecords": 16585,
|
||||||
"totals": {
|
"totals": {
|
||||||
"attributableFailureCount": 5909,
|
"attributableFailureCount": 5910,
|
||||||
"decisionFailureRate": 0.862,
|
"decisionFailureRate": 0.862,
|
||||||
"decisionSuccessRate": 0.138,
|
"decisionSuccessRate": 0.138,
|
||||||
"decisionTotal": 6855,
|
"decisionTotal": 6856,
|
||||||
"failureBreakdown": {
|
"failureBreakdown": {
|
||||||
"ambiguous_runtime": 4024,
|
"ambiguous_runtime": 4026,
|
||||||
"architecture_compatibility": 212,
|
"architecture_compatibility": 212,
|
||||||
"attention_backend": 1,
|
"attention_backend": 1,
|
||||||
"backend_operator": 104,
|
"backend_operator": 104,
|
||||||
"context_length": 321,
|
"context_length": 321,
|
||||||
"framework_architecture_unsupported": 2087,
|
"framework_architecture_unsupported": 2088,
|
||||||
"memory_capacity": 1199,
|
"memory_capacity": 1199,
|
||||||
"model_load": 499,
|
"model_load": 499,
|
||||||
"platform_infrastructure": 925,
|
"platform_infrastructure": 925,
|
||||||
@@ -37661,31 +37727,31 @@
|
|||||||
"日志缺失": 719,
|
"日志缺失": 719,
|
||||||
"验证失败": 674
|
"验证失败": 674
|
||||||
},
|
},
|
||||||
"failureCount": 15442,
|
"failureCount": 15445,
|
||||||
"failureRate": 0.9423,
|
"failureRate": 0.9423,
|
||||||
"pendingCount": 0,
|
"pendingCount": 0,
|
||||||
"pendingRate": 0.0,
|
"pendingRate": 0.0,
|
||||||
"platformFailureCount": 925,
|
"platformFailureCount": 925,
|
||||||
"successCount": 946,
|
"successCount": 946,
|
||||||
"successRate": 0.0577,
|
"successRate": 0.0577,
|
||||||
"total": 16388,
|
"total": 16391,
|
||||||
"unresolvedFailureCount": 8608
|
"unresolvedFailureCount": 8610
|
||||||
},
|
},
|
||||||
"warnings": [
|
"warnings": [
|
||||||
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Iluvatar_mrv-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Mthreads_s4000 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
|
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Kunlunxin_p-800 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Kunlunxin_p-800 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
|
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Biren_166m 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Iluvatar_bi-100 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Sunrise_pt-200-x1 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Sunrise_pt-200-x1 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
|
||||||
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Iluvatar_bi-150 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高(≥50%),建议重点关注。",
|
|
||||||
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU hygon_k100-ai 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Ascend_910-b4 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"GPU MetaX_c-500 本地统计失败率偏高(≥50%),建议重点关注。",
|
"GPU Vastai_va16 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||||
"组合 Iluvatar_bi-150|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Iluvatar_bi-150|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||||
@@ -37725,6 +37791,6 @@
|
|||||||
]
|
]
|
||||||
},
|
},
|
||||||
"storageMode": "decision_state_only",
|
"storageMode": "decision_state_only",
|
||||||
"summarizedRecords": 16582,
|
"summarizedRecords": 16585,
|
||||||
"version": 1
|
"version": 1
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -14,13 +14,13 @@
|
|||||||
100
|
100
|
||||||
],
|
],
|
||||||
"accounts": 12,
|
"accounts": 12,
|
||||||
"activeScanned": 1163,
|
"activeScanned": 1198,
|
||||||
"ageCleanupMode": "admission_only",
|
"ageCleanupMode": "admission_only",
|
||||||
"agePolicySkipped": {
|
"agePolicySkipped": {
|
||||||
"cleanupDisabled": true,
|
"cleanupDisabled": true,
|
||||||
"reason": "admission_only"
|
"reason": "admission_only"
|
||||||
},
|
},
|
||||||
"architectureBlockCount": 139,
|
"architectureBlockCount": 140,
|
||||||
"architectureFrameworkCatalog": {
|
"architectureFrameworkCatalog": {
|
||||||
"ascend_910-b3|text-generation": [
|
"ascend_910-b3|text-generation": [
|
||||||
"llamacpp",
|
"llamacpp",
|
||||||
@@ -98,7 +98,7 @@
|
|||||||
"frameworkCatalogUnknown": 0,
|
"frameworkCatalogUnknown": 0,
|
||||||
"frameworkContextUnknown": 63,
|
"frameworkContextUnknown": 63,
|
||||||
"modelArchitectureUnknown": 133,
|
"modelArchitectureUnknown": 133,
|
||||||
"noMatchingBlock": 1010,
|
"noMatchingBlock": 1045,
|
||||||
"partiallyBlockedFrameworkSet": 20,
|
"partiallyBlockedFrameworkSet": 20,
|
||||||
"runningMatchedProtected": 0,
|
"runningMatchedProtected": 0,
|
||||||
"submissionContextMismatch": 0,
|
"submissionContextMismatch": 0,
|
||||||
@@ -146,5 +146,5 @@
|
|||||||
"repositorySizeUnknown": 0
|
"repositorySizeUnknown": 0
|
||||||
},
|
},
|
||||||
"stopErrors": [],
|
"stopErrors": [],
|
||||||
"uniqueModels": 384
|
"uniqueModels": 383
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -217,6 +217,7 @@
|
|||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T13:30:42.155984+00:00", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23524076260, "estimatedRequiredGiB": 26.327, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 23556606333, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6649116524, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 23556606333}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T05:16:13.136785+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4971300", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T13:30:42.155984+00:00", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23524076260, "estimatedRequiredGiB": 26.327, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 23556606333, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6649116524, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 23556606333}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T05:16:13.136785+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4971300", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-21T11:15:42.052921+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T05:00:40.870810+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4971125", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-21T11:15:42.052921+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T05:00:40.870810+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4971125", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "参数/模板问题", "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T20:49:19.683254+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T04:46:07.939821+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970952", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": "参数/模板问题", "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T20:49:19.683254+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T04:46:07.939821+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970952", "taskType": "text-generation", "verifyResult": null}
|
||||||
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-21T19:57:11.769116+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:44:53.386157+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970949", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["mistral3"], "framework": "vllm", "lastSyncTime": "2026-09-21T04:36:00.569168+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:42:19.140494+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970923", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["mistral3"], "framework": "vllm", "lastSyncTime": "2026-09-21T04:36:00.569168+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:42:19.140494+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970923", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:58:28.353121+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:42:19.137834+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970925", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:58:28.353121+00:00", "modelId": "mlx-community/gemma-4-e4b-it-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7486327132, "estimatedRequiredGiB": 8.403, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 7518908360, "modelscopeLicense": "gemma", "modelscopeParams": 2227418442, "modelscopeTags": ["license:gemma", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7518908360}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:42:19.137834+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970925", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4_unified"], "framework": "vllm", "lastSyncTime": "2026-09-21T15:02:24.876330+00:00", "modelId": "mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9003966484, "estimatedRequiredGiB": 10.099, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 9036402901, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2463563824, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9036402901}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:42:19.135971+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970924", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4_unified"], "framework": "vllm", "lastSyncTime": "2026-09-21T15:02:24.876330+00:00", "modelId": "mlx-community/gemma-4-12B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9003966484, "estimatedRequiredGiB": 10.099, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 9036402901, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2463563824, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9036402901}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T04:42:19.135971+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970924", "taskType": "text-generation", "verifyResult": -1}
|
||||||
@@ -276,6 +277,7 @@
|
|||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:21:49.902144+00:00", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23524076260, "estimatedRequiredGiB": 26.327, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 23556606333, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6649116524, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 23556606333}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.460457+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970350", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:21:49.902144+00:00", "modelId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 23524076260, "estimatedRequiredGiB": 26.327, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 23556606333, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6649116524, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 23556606333}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.460457+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970350", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-21T01:35:45.556133+00:00", "modelId": "JANGQ-AI/AppleScript-8B-JANG_4M", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T03:58:35.458351+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970351", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": "参数/模板问题", "framework": "vllm", "lastSyncTime": "2026-09-21T01:35:45.556133+00:00", "modelId": "JANGQ-AI/AppleScript-8B-JANG_4M", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T03:58:35.458351+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970351", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:21:49.902097+00:00", "modelId": "XHToken/Spark-X2.5-4B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8224192408, "estimatedRequiredGiB": 9.209, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 8239718268, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4112079360, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8239718268}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.457045+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970357", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:21:49.902097+00:00", "modelId": "XHToken/Spark-X2.5-4B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8224192408, "estimatedRequiredGiB": 9.209, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 8239718268, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4112079360, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8239718268}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.457045+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970357", "taskType": "text-generation", "verifyResult": -1}
|
||||||
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_vl"], "framework": "vllm", "lastSyncTime": "2026-09-21T19:57:11.769130+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.455197+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970358", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:21:49.901805+00:00", "modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21855018733, "estimatedRequiredGiB": 24.461, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 21887495988, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6144703054, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:moe", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21887495988}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.451794+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970348", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["gemma4"], "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:21:49.901805+00:00", "modelId": "mlx-community/gemma-4-26B-A4B-it-qat-OptiQ-4bit", "modelProfile": {"architectures": ["Gemma4ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 21855018733, "estimatedRequiredGiB": 24.461, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4", "modelscopeFileSize": 21887495988, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6144703054, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:qat", "custom_tag:moe", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:image-text-to-text", "custom_tag:gemma-4"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 21887495988}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.451794+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970348", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "memory_capacity", "failureAction": "reject_if_estimated_load_exceeds_memory", "failureCategory": "memory_capacity", "failureClassificationReason": "structured_oom", "failureCode": "PREFLIGHT_OOM", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureObservedGpuMemoryGiB": 32.0, "failureScope": "model_gpu", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:21:49.901895+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683241105, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703487237, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703487237}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.449495+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970352", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "memory_capacity", "failureAction": "reject_if_estimated_load_exceeds_memory", "failureCategory": "memory_capacity", "failureClassificationReason": "structured_oom", "failureCode": "PREFLIGHT_OOM", "failureDetectedFramework": "vllm_fix_tokenizer", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureObservedGpuMemoryGiB": 32.0, "failureScope": "model_gpu", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:21:49.901895+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20683241105, "estimatedRequiredGiB": 23.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20703487237, "modelscopeLicense": "apache-2.0", "modelscopeParams": 6086364400, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:macos", "custom_tag:speculative-decoding", "custom_tag:multi-token-prediction", "custom_tag:qwen", "custom_tag:qwen3.8", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:coding", "custom_tag:qwen3-8", "custom_tag:qwen-3.8", "custom_tag:local-llm", "custom_tag:llm", "custom_tag:m5", "custom_tag:m5-max", "custom_tag:m4", "custom_tag:m3", "custom_tag:macbook-pro", "custom_tag:mac-studio", "custom_tag:opencode", "custom_tag:claude-code", "custom_tag:27b", "custom_tag:qwen3.8-27b", "custom_tag:qwen3-8-27b"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 20703487237}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.449495+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970352", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T19:49:37.657161+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.441724+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970354", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T19:49:37.657161+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:58:35.441724+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970354", "taskType": "text-generation", "verifyResult": -1}
|
||||||
@@ -296,5 +298,3 @@
|
|||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T15:02:24.876345+00:00", "modelId": "neuralmagic/starcoder2-15b-FP8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16564748304, "estimatedRequiredGiB": 18.516, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 16568142341, "modelscopeLicense": "other", "modelscopeParams": 15957889024, "modelscopeTags": ["license:other", "model_type:starcoder2", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 16568142341}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:31.484685+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970312", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T15:02:24.876345+00:00", "modelId": "neuralmagic/starcoder2-15b-FP8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16564748304, "estimatedRequiredGiB": 18.516, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 16568142341, "modelscopeLicense": "other", "modelscopeParams": 15957889024, "modelscopeTags": ["license:other", "model_type:starcoder2", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 16568142341}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:57:31.484685+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970312", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "repository_structure", "failureAction": "require_framework_specific_root_files", "failureCategory": "repository_structure", "failureClassificationReason": "structured_missing_files", "failureCode": "MODEL_FILE_NOT_FOUND", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T16:21:18.661407+00:00", "modelId": "aisingapore/SEA-LION-v1-7B-IT-Research", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "ecd31d4ed4e0e92b9deabe3de0f18c75dc252004ba68826e4cfd561549549074", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008118474, "modelscopeLicense": "cc-by-nc-sa-4.0", "modelscopeParams": 7501651968, "modelscopeTags": ["license:cc-by-nc-sa-4.0", "model_type:mpt", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008118474}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:53:33.643969+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970234", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "repository_structure", "failureAction": "require_framework_specific_root_files", "failureCategory": "repository_structure", "failureClassificationReason": "structured_missing_files", "failureCode": "MODEL_FILE_NOT_FOUND", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T16:21:18.661407+00:00", "modelId": "aisingapore/SEA-LION-v1-7B-IT-Research", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "ecd31d4ed4e0e92b9deabe3de0f18c75dc252004ba68826e4cfd561549549074", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008118474, "modelscopeLicense": "cc-by-nc-sa-4.0", "modelscopeParams": 7501651968, "modelscopeTags": ["license:cc-by-nc-sa-4.0", "model_type:mpt", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008118474}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:53:33.643969+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970234", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T22:43:47.754882+00:00", "modelId": "IntervitensInc/kek_mk3", "modelProfile": {"architectures": ["StableLmForCausalLM"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3289069288, "estimatedRequiredGiB": 3.683, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "stablelm", "modelscopeFileSize": 3295875324, "modelscopeLicense": null, "modelscopeParams": 1644515328, "modelscopeTags": ["model_type:stablelm", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3295875324}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.813804+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970120", "taskType": "text-generation", "verifyResult": -1}
|
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T22:43:47.754882+00:00", "modelId": "IntervitensInc/kek_mk3", "modelProfile": {"architectures": ["StableLmForCausalLM"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3289069288, "estimatedRequiredGiB": 3.683, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "stablelm", "modelscopeFileSize": 3295875324, "modelscopeLicense": null, "modelscopeParams": 1644515328, "modelscopeTags": ["model_type:stablelm", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3295875324}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.813804+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970120", "taskType": "text-generation", "verifyResult": -1}
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "missing_structured_error_code", "failureCode": null, "failureDetectedFramework": "transformers", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "transformers", "lastSyncTime": "2026-09-20T14:14:39.678925+00:00", "modelId": "aisingapore/SEA-LION-v1-7B", "modelProfile": {"architectures": ["MPTForCausalLM"], "configFingerprint": "0335c7e4668aaa963303ebe91f80d6eb1e7e8d0299011bd85f7e05a4df031dc7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15003361968, "estimatedRequiredGiB": 16.773, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mpt", "modelscopeFileSize": 15008115847, "modelscopeLicense": "mit", "modelscopeParams": 7501651968, "modelscopeTags": ["license:mit", "model_type:mpt", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 15008115847}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.801267+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970119", "taskType": "text-generation", "verifyResult": -1}
|
|
||||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T17:09:45.266093+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29274355224, "estimatedRequiredGiB": 32.761, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 29314389092, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 29314389092}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:46:35.789758+00:00", "targetGpu": "Biren_166m", "taskId": "4970118", "taskType": "text-generation", "verifyResult": -1}
|
|
||||||
|
|||||||
@@ -2286,6 +2286,8 @@
|
|||||||
{"batchId": "af01956827ff4c169cbafc6bc6e203fe", "completedAt": "2026-09-21T19:52:49.937981+00:00", "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:52:48.876602+00:00", "framework": "vllm", "intentId": "27db0d5b0ed2446a9087527b204e962b", "lastModified": "2026-09-21T18:48:50+00:00", "modelAddress": "https://modelscope.cn/models/prithivMLmods/Nexus-9B-CodeCore-Merge", "reason": null, "reconciledAt": "2026-09-21T19:54:21.694882+00:00", "repoId": "prithivMLmods/Nexus-9B-CodeCore-Merge", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Iluvatar_mrv-100", "taskId": "5006513", "taskType": "text-generation"}
|
{"batchId": "af01956827ff4c169cbafc6bc6e203fe", "completedAt": "2026-09-21T19:52:49.937981+00:00", "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:52:48.876602+00:00", "framework": "vllm", "intentId": "27db0d5b0ed2446a9087527b204e962b", "lastModified": "2026-09-21T18:48:50+00:00", "modelAddress": "https://modelscope.cn/models/prithivMLmods/Nexus-9B-CodeCore-Merge", "reason": null, "reconciledAt": "2026-09-21T19:54:21.694882+00:00", "repoId": "prithivMLmods/Nexus-9B-CodeCore-Merge", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Iluvatar_mrv-100", "taskId": "5006513", "taskType": "text-generation"}
|
||||||
{"batchId": "af01956827ff4c169cbafc6bc6e203fe", "completedAt": "2026-09-21T19:52:49.937983+00:00", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:52:48.876654+00:00", "framework": "vllm", "intentId": "a805e242301145ed90037bde1881f57f", "lastModified": "2026-09-21T18:24:38+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-attnres-block2-v1_B", "reason": null, "reconciledAt": "2026-09-21T19:54:21.693692+00:00", "repoId": "nkkbr/Mini-K3-1H-attnres-block2-v1_B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Mthreads_s4000", "taskId": "5006517", "taskType": "text-generation"}
|
{"batchId": "af01956827ff4c169cbafc6bc6e203fe", "completedAt": "2026-09-21T19:52:49.937983+00:00", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:52:48.876654+00:00", "framework": "vllm", "intentId": "a805e242301145ed90037bde1881f57f", "lastModified": "2026-09-21T18:24:38+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-attnres-block2-v1_B", "reason": null, "reconciledAt": "2026-09-21T19:54:21.693692+00:00", "repoId": "nkkbr/Mini-K3-1H-attnres-block2-v1_B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Mthreads_s4000", "taskId": "5006517", "taskType": "text-generation"}
|
||||||
{"batchId": "af01956827ff4c169cbafc6bc6e203fe", "completedAt": "2026-09-21T19:52:49.937985+00:00", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:52:48.876728+00:00", "framework": "vllm", "intentId": "e77e8806b18040e1be44cfa579fd0eed", "lastModified": "2026-09-21T18:27:54+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-decay-g1-v2_B", "reason": null, "reconciledAt": "2026-09-21T19:54:21.693963+00:00", "repoId": "nkkbr/Mini-K3-1H-decay-g1-v2_B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Mthreads_s4000", "taskId": "5006515", "taskType": "text-generation"}
|
{"batchId": "af01956827ff4c169cbafc6bc6e203fe", "completedAt": "2026-09-21T19:52:49.937985+00:00", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:52:48.876728+00:00", "framework": "vllm", "intentId": "e77e8806b18040e1be44cfa579fd0eed", "lastModified": "2026-09-21T18:27:54+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-decay-g1-v2_B", "reason": null, "reconciledAt": "2026-09-21T19:54:21.693963+00:00", "repoId": "nkkbr/Mini-K3-1H-decay-g1-v2_B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Mthreads_s4000", "taskId": "5006515", "taskType": "text-generation"}
|
||||||
|
{"batchId": "f3bf040e0b724f75a3a10b60f89f5ff7", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:57:32.018229+00:00", "framework": "vllm", "intentId": "7372e9b97294486d8562e867a0c2000d", "lastModified": "2026-09-21T18:34:41+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-32-v2_D", "repoId": "nkkbr/Mini-K3-1H-kda-kernel-32-v2_D", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Mthreads_s4000", "taskType": "text-generation"}
|
||||||
|
{"batchId": "f3bf040e0b724f75a3a10b60f89f5ff7", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:57:32.018364+00:00", "framework": "vllm", "intentId": "a04a2eacdcd84031a868f37216f583a8", "lastModified": "2026-09-21T18:33:37+00:00", "modelAddress": "https://modelscope.cn/models/nkkbr/Mini-K3-1H-kda-kernel-8-v2_D", "repoId": "nkkbr/Mini-K3-1H-kda-kernel-8-v2_D", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Mthreads_s4000", "taskType": "text-generation"}
|
||||||
{"batchId": "f8d3f1991720478ea905a5ecf7d43601", "completedAt": "2026-09-21T19:05:00.691781+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:00:55.137064+00:00", "framework": "vllm", "intentId": "64b834dd640742a0b32535bde02b94ad", "repoId": "CohereLabs/tiny-aya-base-32K", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "age_policy_deferred", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
|
{"batchId": "f8d3f1991720478ea905a5ecf7d43601", "completedAt": "2026-09-21T19:05:00.691781+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:00:55.137064+00:00", "framework": "vllm", "intentId": "64b834dd640742a0b32535bde02b94ad", "repoId": "CohereLabs/tiny-aya-base-32K", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "age_policy_deferred", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
|
||||||
{"batchId": "f8d3f1991720478ea905a5ecf7d43601", "completedAt": "2026-09-21T19:05:00.691779+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:00:55.137007+00:00", "framework": "vllm", "intentId": "d6229ef4ac9942078baa35e312133d9d", "repoId": "neuralmagic/Llama-3.2-1B-Instruct-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "age_policy_deferred", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
|
{"batchId": "f8d3f1991720478ea905a5ecf7d43601", "completedAt": "2026-09-21T19:05:00.691779+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:00:55.137007+00:00", "framework": "vllm", "intentId": "d6229ef4ac9942078baa35e312133d9d", "repoId": "neuralmagic/Llama-3.2-1B-Instruct-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "age_policy_deferred", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
|
||||||
{"batchId": "f8d3f1991720478ea905a5ecf7d43601", "completedAt": "2026-09-21T19:05:00.691776+00:00", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:00:55.136949+00:00", "framework": "vllm", "intentId": "cff6b666a842431faf928f94c2086fdf", "repoId": "nm-testing/llama-3-instruct-w8a8-dyn-per-token-test", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Mthreads_s4000", "taskId": null, "taskType": "text-generation"}
|
{"batchId": "f8d3f1991720478ea905a5ecf7d43601", "completedAt": "2026-09-21T19:05:00.691776+00:00", "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configSource": "modelhub_live", "createdAt": "2026-09-21T19:00:55.136949+00:00", "framework": "vllm", "intentId": "cff6b666a842431faf928f94c2086fdf", "repoId": "nm-testing/llama-3-instruct-w8a8-dyn-per-token-test", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "age_policy_deferred", "targetGpu": "Mthreads_s4000", "taskId": null, "taskType": "text-generation"}
|
||||||
|
|||||||
@@ -2,24 +2,24 @@
|
|||||||
"agentVersion": "2026.09.20.2",
|
"agentVersion": "2026.09.20.2",
|
||||||
"checksums": {
|
"checksums": {
|
||||||
".modelhub_state/account_capacity.json": "68d2a8390ad022f2ddcd30841f3860c4535387fa64cb46dc7507eb16837e798f",
|
".modelhub_state/account_capacity.json": "68d2a8390ad022f2ddcd30841f3860c4535387fa64cb46dc7507eb16837e798f",
|
||||||
".modelhub_state/architecture_compatibility_blacklist.json": "f5f85558525a4d76195e4f38be00f338191304994d8625d3f0847492aa0a09e4",
|
".modelhub_state/architecture_compatibility_blacklist.json": "d9d3d3af7b4468bbb39e022fd3b55d57cc8de5a95f0ad6ab3c00b59fd15909b0",
|
||||||
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
||||||
".modelhub_state/market_intelligence.json": "807724d8f9f15d683ef21c858b620a5d7a6ee0e02167b524f0f888787ae2cfd6",
|
".modelhub_state/market_intelligence.json": "d64aec77faa150e08fd405453535fe01984f983049e10ec2172bd8f0a715b7ac",
|
||||||
".modelhub_state/official_capabilities.json": "ae378ddb1825eb575792cf41ed3fa445d70a656b7bda7f9fa8c02babe3b6b875",
|
".modelhub_state/official_capabilities.json": "d693ea775df4b7b618c3483b5c4283aecbd103c403ad3a87ae8e22401634d9e3",
|
||||||
".modelhub_state/outcome_checkpoint.json": "f31ae2c0f5bf9f13e351483e00ddd81b10ef4e16337af0ff95802bba684c82fd",
|
".modelhub_state/outcome_checkpoint.json": "d77e18ebd210cc8796fde5a337f60644c50770f7cd8d49c0706d987e5ba6ff7e",
|
||||||
".modelhub_state/queue_cleanup_latest.json": "fab470f2bbf8c39d73c8917d9fdab78c775c47a597aa2530a529e5a7ae55d4ec",
|
".modelhub_state/queue_cleanup_latest.json": "f986d7b722dd04a2b60ce98748799c4497653d40e223f8abf91fb22a8aa6989e",
|
||||||
".modelhub_state/recent_outcomes.jsonl": "0e8003011fc7ffe4b589f4c251a9a27dcc4a692e1ded9860e25289e730723bf3",
|
".modelhub_state/recent_outcomes.jsonl": "811ac1cb2df8576c601ab88809bc03272ba71597a3bb6976f9643773cb3f6dcc",
|
||||||
".modelhub_state/recovery_active_tasks.jsonl": "a9ffecc7009a682344d7ecabdc9261681080d3f5503b12593d337d3bc4dbfc89",
|
".modelhub_state/recovery_active_tasks.jsonl": "a9ffecc7009a682344d7ecabdc9261681080d3f5503b12593d337d3bc4dbfc89",
|
||||||
".modelhub_state/recovery_intents.jsonl": "b485fbb0cf466960fc30d4466b7ed5bf727c8ec4a8eea61ab8c09b9ecf883f7a",
|
".modelhub_state/recovery_intents.jsonl": "b98d76c6baf1f76ed85272f193597e6393d905a3b4e7a920f5ef9bf8c7033f2e",
|
||||||
".modelhub_state/routing_intelligence.json": "19eb4ba03b4bb3066fb87808c9c10b326fdcd60ea0503432eb1e767add61a40d",
|
".modelhub_state/routing_intelligence.json": "19eb4ba03b4bb3066fb87808c9c10b326fdcd60ea0503432eb1e767add61a40d",
|
||||||
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
||||||
".modelhub_state/worker_crashes.jsonl": "07505b69e0b050da5f90f8b046a3069d3e1ca2574ca39896bc7d8a66293167f7",
|
".modelhub_state/worker_crashes.jsonl": "07505b69e0b050da5f90f8b046a3069d3e1ca2574ca39896bc7d8a66293167f7",
|
||||||
"ledger/submissions.jsonl": "f0fe8171536cea0ffb77827c281571ccdc93c69a86755342ee6383c3ecce1548",
|
"ledger/submissions.jsonl": "f0fe8171536cea0ffb77827c281571ccdc93c69a86755342ee6383c3ecce1548",
|
||||||
"outcomes/submissions.jsonl": "9344e1ea638405bbaefdc6711e10bffbb93486c035ef7e50d0e0b36c46d6a54a"
|
"outcomes/submissions.jsonl": "361156073c270cd2acf69a2d48342fd51d92d4b64e9419719152cebe87028fa8"
|
||||||
},
|
},
|
||||||
"generation": 11675,
|
"generation": 11676,
|
||||||
"phase": "cycle",
|
"phase": "intent",
|
||||||
"schemaVersion": 1,
|
"schemaVersion": 1,
|
||||||
"updatedAt": "2026-09-21T19:56:10.659187+00:00",
|
"updatedAt": "2026-09-21T19:57:32.193910+00:00",
|
||||||
"writerId": "8b35139af6674067a339a670222d4b67"
|
"writerId": "8b35139af6674067a339a670222d4b67"
|
||||||
}
|
}
|
||||||
|
|||||||
@@ -469,7 +469,6 @@
|
|||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.670706+00:00", "modelId": "RedHatAI/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663591, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663591}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:40:44.501444+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969969", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.670706+00:00", "modelId": "RedHatAI/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663591, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663591}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:40:44.501444+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969969", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.671081+00:00", "modelId": "RedHatAI/SmolLM-360M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "f49713b9fe8e3e7683ab2722544c85a44993bcd72b9421c1599f634e68b3a2f7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 504052632, "estimatedRequiredGiB": 0.567, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 507443052, "modelscopeLicense": "apache-2.0", "modelscopeParams": 409007040, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 507443052}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:40:44.505215+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969967", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.671081+00:00", "modelId": "RedHatAI/SmolLM-360M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "f49713b9fe8e3e7683ab2722544c85a44993bcd72b9421c1599f634e68b3a2f7", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 504052632, "estimatedRequiredGiB": 0.567, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 507443052, "modelscopeLicense": "apache-2.0", "modelscopeParams": 409007040, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 507443052}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:40:44.505215+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969967", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.671015+00:00", "modelId": "LiquidAI/LFM2.5-2.6B", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5394427456, "estimatedRequiredGiB": 6.049, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 5412393192, "modelscopeLicense": "other", "modelscopeParams": 2697198592, "modelscopeTags": ["license:other", "model_type:lfm2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5412393192}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:40:44.507984+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969968", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.671015+00:00", "modelId": "LiquidAI/LFM2.5-2.6B", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5394427456, "estimatedRequiredGiB": 6.049, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 5412393192, "modelscopeLicense": "other", "modelscopeParams": 2697198592, "modelscopeTags": ["license:other", "model_type:lfm2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5412393192}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:40:44.507984+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4969968", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:01:00.670523+00:00", "modelId": "neuralmagic/Qwen2-0.5B-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 903168128, "estimatedRequiredGiB": 1.022, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 914658310, "modelscopeLicense": "apache-2.0", "modelscopeParams": 630167424, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 914658310}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:42.237083+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969989", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.671348+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248516007, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248516007}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:42.249155+00:00", "targetGpu": "Vastai_va16", "taskId": "4969994", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.671348+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248516007, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248516007}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:42.249155+00:00", "targetGpu": "Vastai_va16", "taskId": "4969994", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.670941+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-bf16", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2340697867, "estimatedRequiredGiB": 2.621, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 2345554722, "modelscopeLicense": "other", "modelscopeParams": 1170340608, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2345554722}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:42.236412+00:00", "targetGpu": "Vastai_va16", "taskId": "4969990", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.670941+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-bf16", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2340697867, "estimatedRequiredGiB": 2.621, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 2345554722, "modelscopeLicense": "other", "modelscopeParams": 1170340608, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2345554722}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:42.236412+00:00", "targetGpu": "Vastai_va16", "taskId": "4969990", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.671225+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:42.238130+00:00", "targetGpu": "Vastai_va16", "taskId": "4969985", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.671225+00:00", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955963132, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx", "custom_tag:reasoning"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955963132}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:41:42.238130+00:00", "targetGpu": "Vastai_va16", "taskId": "4969985", "taskType": "text-generation", "verifyResult": null}
|
||||||
@@ -548,7 +547,6 @@
|
|||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.671026+00:00", "modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8674988799, "estimatedRequiredGiB": 9.718, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8695119291, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2415484144, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:speculative-decoding", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:qwen3-5", "custom_tag:9b", "custom_tag:6-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8695119291}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:35.439990+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970355", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:01:00.671026+00:00", "modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8674988799, "estimatedRequiredGiB": 9.718, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8695119291, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2415484144, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:apple-silicon", "custom_tag:speculative-decoding", "custom_tag:qwen", "custom_tag:qwen3", "custom_tag:mtp", "custom_tag:mtplx", "custom_tag:local-ai", "custom_tag:qwen3-5", "custom_tag:9b", "custom_tag:6-bit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8695119291}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:35.439990+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970355", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:01:00.670624+00:00", "modelId": "RedHatAI/SmolLM-135M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "26bfee231695e882b96dee0191bc1dfec25bcad132af96a7ee5f0e26f3cb9fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 219850592, "estimatedRequiredGiB": 0.249, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 223236688, "modelscopeLicense": "apache-2.0", "modelscopeParams": 162826560, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 223236688}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:35.535522+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970356", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:01:00.670624+00:00", "modelId": "RedHatAI/SmolLM-135M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "26bfee231695e882b96dee0191bc1dfec25bcad132af96a7ee5f0e26f3cb9fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 219850592, "estimatedRequiredGiB": 0.249, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 223236688, "modelscopeLicense": "apache-2.0", "modelscopeParams": 162826560, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 223236688}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:35.535522+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970356", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:01:00.671236+00:00", "modelId": "RedHatAI/gemma-2-27b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 30775446656, "estimatedRequiredGiB": 34.419, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 30797371689, "modelscopeLicense": "gemma", "modelscopeParams": 28406776320, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 30797371689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:35.447710+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970349", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:01:00.671236+00:00", "modelId": "RedHatAI/gemma-2-27b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 30775446656, "estimatedRequiredGiB": 34.419, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 30797371689, "modelscopeLicense": "gemma", "modelscopeParams": 28406776320, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 30797371689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:35.447710+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970349", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.670884+00:00", "modelId": "BAAI/RoboBrain2.5-8B-MT", "modelProfile": {"architectures": ["Qwen3VLForConditionalGeneration"], "configFingerprint": "6db01b5037d5ed2db74875d2113f8582c97f34bd825637161e54174c9f2cabed", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17534339552, "estimatedRequiredGiB": 19.609, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_vl", "modelscopeFileSize": 17545918174, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8767123696, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_vl", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17545918174}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:35.455197+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970358", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.671020+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:41.296912+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4970360", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:01:00.671020+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:58:41.296912+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4970360", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:04:34.964276+00:00", "modelId": "neuralmagic/gemma-2-2b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4385544296, "estimatedRequiredGiB": 4.926, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 4407367987, "modelscopeLicense": "gemma", "modelscopeParams": 3204165888, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4407367987}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:01:42.646836+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4970409", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:04:34.964276+00:00", "modelId": "neuralmagic/gemma-2-2b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4385544296, "estimatedRequiredGiB": 4.926, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 4407367987, "modelscopeLicense": "gemma", "modelscopeParams": 3204165888, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4407367987}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:01:42.646836+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4970409", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:04:34.964288+00:00", "modelId": "neuralmagic/Meta-Llama-3.1-8B-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084040952, "estimatedRequiredGiB": 10.163, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093251901, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093251901}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:02:45.317909+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4970439", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:04:34.964288+00:00", "modelId": "neuralmagic/Meta-Llama-3.1-8B-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9084040952, "estimatedRequiredGiB": 10.163, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 9093251901, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 9093251901}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:02:45.317909+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4970439", "taskType": "text-generation", "verifyResult": null}
|
||||||
@@ -593,7 +591,6 @@
|
|||||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T12:34:29.856347+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:31:05.659043+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970789", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T12:34:29.856347+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:31:05.659043+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970789", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:46:14.462961+00:00", "modelId": "RedHatAI/gemma-2-2b-it-quantized.w4a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3403624536, "estimatedRequiredGiB": 3.828, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 3425449433, "modelscopeLicense": "llama2", "modelscopeParams": 3204165888, "modelscopeTags": ["license:llama2", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 3425449433}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:36:48.898217+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4970869", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:46:14.462961+00:00", "modelId": "RedHatAI/gemma-2-2b-it-quantized.w4a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3403624536, "estimatedRequiredGiB": 3.828, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 3425449433, "modelscopeLicense": "llama2", "modelscopeParams": 3204165888, "modelscopeTags": ["license:llama2", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 3425449433}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:36:48.898217+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4970869", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:46:14.462938+00:00", "modelId": "empero-ai/Qwen3.8-35B-A3B-Distill", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 71903778536, "estimatedRequiredGiB": 80.392, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 71933964484, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.6", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:moe", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 2}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-15T22:59:41.110897+00:00", "modelType": "qwen3_5_moe", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 2}, "repositoryOnDiskBytes": 71933964484}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:42:19.067096+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970926", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:46:14.462938+00:00", "modelId": "empero-ai/Qwen3.8-35B-A3B-Distill", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 71903778536, "estimatedRequiredGiB": 80.392, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 71933964484, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:empero-ai", "custom_tag:qwen3.6", "custom_tag:qwen3.8", "custom_tag:distillation", "custom_tag:reasoning", "custom_tag:function-calling", "custom_tag:moe", "custom_tag:sft"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 2}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-15T22:59:41.110897+00:00", "modelType": "qwen3_5_moe", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 2}, "repositoryOnDiskBytes": 71933964484}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:42:19.067096+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970926", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-20T12:46:14.462954+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 951093016, "estimatedRequiredGiB": 1.068, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 955857175, "modelscopeLicense": "other", "modelscopeParams": 256113408, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 955857175}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:44:53.386157+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4970949", "taskType": "text-generation", "verifyResult": null}
|
|
||||||
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:46:14.462918+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248409935, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248409935}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:44:53.390325+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970950", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-20T12:46:14.462918+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248409935, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248409935}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:44:53.390325+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4970950", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T13:05:06.219610+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T05:01:56.033898+00:00", "targetGpu": "Biren_166m", "taskId": "4971137", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T13:05:06.219610+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T05:01:56.033898+00:00", "targetGpu": "Biren_166m", "taskId": "4971137", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T13:18:18.977292+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T05:16:13.324611+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4971311", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T13:18:18.977292+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T05:16:13.324611+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4971311", "taskType": "text-generation", "verifyResult": null}
|
||||||
@@ -1005,7 +1002,7 @@
|
|||||||
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-21T19:49:37.657110+00:00", "modelId": "RedHatAI/starcoder2-3b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4093180600, "estimatedRequiredGiB": 4.578, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 4096479058, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 3181366272, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4096479058}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:46:34.674727+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000249", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-21T19:49:37.657110+00:00", "modelId": "RedHatAI/starcoder2-3b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4093180600, "estimatedRequiredGiB": 4.578, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 4096479058, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 3181366272, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4096479058}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:46:34.674727+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000249", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-21T19:49:37.657141+00:00", "modelId": "RedHatAI/starcoder2-7b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7857298896, "estimatedRequiredGiB": 8.785, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7860673720, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 7400416256, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7860673720}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:48:23.283753+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000251", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm-customized", "lastSyncTime": "2026-09-21T19:49:37.657141+00:00", "modelId": "RedHatAI/starcoder2-7b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7857298896, "estimatedRequiredGiB": 8.785, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7860673720, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 7400416256, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7860673720}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:48:23.283753+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "5000251", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T19:52:34.157576+00:00", "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4395007696, "estimatedRequiredGiB": 4.922, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 4404161088, "modelscopeLicense": "llama3.2", "modelscopeParams": 3606752256, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4404161088}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:50:46.187626+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5000279", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T19:52:34.157576+00:00", "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4395007696, "estimatedRequiredGiB": 4.922, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 4404161088, "modelscopeLicense": "llama3.2", "modelscopeParams": 3606752256, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4404161088}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:50:46.187626+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5000279", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T11:54:31.274098+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000347", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T19:57:11.769093+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Instruct-4bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 658540250, "estimatedRequiredGiB": 0.741, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 663396475, "modelscopeLicense": "other", "modelscopeParams": 182975232, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 663396475}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T11:54:31.274098+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000347", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "RedHatAI/Phi-3-mini-128k-instruct-FP8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4018332584, "estimatedRequiredGiB": 4.494, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4020783717, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4020783717}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T11:57:55.153765+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5000419", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "RedHatAI/Phi-3-mini-128k-instruct-FP8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4018332584, "estimatedRequiredGiB": 4.494, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4020783717, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4020783717}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T11:57:55.153765+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5000419", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29274355224, "estimatedRequiredGiB": 32.761, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 29314389092, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 29314389092}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T11:57:55.143712+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000420", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29274355224, "estimatedRequiredGiB": 32.761, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 29314389092, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 29314389092}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T11:57:55.143712+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000420", "taskType": "text-generation", "verifyResult": null}
|
||||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w4a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7686303632, "estimatedRequiredGiB": 8.592, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 7688299781, "modelscopeLicense": "llama2", "modelscopeParams": 13960238080, "modelscopeTags": ["license:llama2", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7688299781}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T11:57:55.174530+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000425", "taskType": "text-generation", "verifyResult": null}
|
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w4a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7686303632, "estimatedRequiredGiB": 8.592, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 7688299781, "modelscopeLicense": "llama2", "modelscopeParams": 13960238080, "modelscopeTags": ["license:llama2", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7688299781}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T11:57:55.174530+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "5000425", "taskType": "text-generation", "verifyResult": null}
|
||||||
|
|||||||
Reference in New Issue
Block a user