state: generation 5230 (intent)

This commit is contained in:
2026-09-10 15:29:43 +00:00
parent 35f5d18664
commit 4eaabbbd50
9 changed files with 1144 additions and 1145 deletions

View File

@@ -1311,7 +1311,7 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-10T15:23:51.438109+00:00",
"generatedAt": "2026-09-10T15:27:10.037984+00:00",
"summary": {
"activeBlockCount": 66,
"byGpuFramework": {

View File

@@ -3,106 +3,106 @@
"communityError": null,
"communitySample": {},
"communityUpdatedAt": "2026-09-10T15:16:33.850208+00:00",
"frameworkAttemptedAt": "2026-09-10T12:35:14.692612+00:00",
"frameworkAttemptedAt": "2026-09-10T15:27:32.924657+00:00",
"frameworkError": null,
"frameworkStats": {
"text-generation": {
"Ascend_910-b3": {
"llamacpp": {
"framework": "llamacpp",
"modelCount": 38318,
"modelCount": 38338,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 18283,
"successRate": 0.47713868155958034,
"wilsonLowerBound": 0.47214006949864223
"successCount": 18295,
"successRate": 0.47720277531430955,
"wilsonLowerBound": 0.47220543078048205
},
"vllm": {
"framework": "vllm",
"modelCount": 50542,
"modelCount": 50568,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 14324,
"successRate": 0.28340785881049424,
"wilsonLowerBound": 0.27949552744451256
"successCount": 14328,
"successRate": 0.28334124347413386,
"wilsonLowerBound": 0.27943019786424367
},
"vllm_tokenizer_patch": {
"framework": "vllm_tokenizer_patch",
"modelCount": 1770,
"modelCount": 1789,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 445,
"successRate": 0.2514124293785311,
"wilsonLowerBound": 0.2317546860465486
"successCount": 449,
"successRate": 0.2509782001117943,
"wilsonLowerBound": 0.23143456102944407
}
},
"Ascend_910-b4": {
"llamacpp": {
"framework": "llamacpp",
"modelCount": 31396,
"modelCount": 31416,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 13506,
"successRate": 0.43018218881386167,
"wilsonLowerBound": 0.42471443238148554
"successCount": 13521,
"successRate": 0.43038579067990834,
"wilsonLowerBound": 0.42491943017841394
},
"vllm": {
"framework": "vllm",
"modelCount": 36514,
"modelCount": 36525,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 10437,
"successRate": 0.2858355699183875,
"wilsonLowerBound": 0.2812239942679824
"successCount": 10442,
"successRate": 0.285886379192334,
"wilsonLowerBound": 0.28127524230104706
},
"vllm_tokenizer_patch": {
"framework": "vllm_tokenizer_patch",
"modelCount": 6942,
"modelCount": 6968,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3194,
"successRate": 0.460097954479977,
"wilsonLowerBound": 0.4483986894698718
"successCount": 3199,
"successRate": 0.4590987370838117,
"wilsonLowerBound": 0.4474237173995597
}
},
"Biren_166m": {
"vllm": {
"framework": "vllm",
"modelCount": 63790,
"modelCount": 63793,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 12060,
"successRate": 0.18905784605737577,
"wilsonLowerBound": 0.1860380147747332
"successRate": 0.18904895521452197,
"wilsonLowerBound": 0.18602924981833124
},
"vllm_fix_tokenizer": {
"framework": "vllm_fix_tokenizer",
"modelCount": 13800,
"modelCount": 13844,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 2726,
"successRate": 0.19753623188405797,
"wilsonLowerBound": 0.19097797620808124
"successCount": 2735,
"successRate": 0.19755850910141579,
"wilsonLowerBound": 0.1910102609114936
},
"vllm_tokenizer_patch": {
"framework": "vllm_tokenizer_patch",
"modelCount": 5147,
"modelCount": 5151,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 36,
"successRate": 0.006994365649893142,
"wilsonLowerBound": 0.005056576682829683
"successCount": 37,
"successRate": 0.007183071248301301,
"wilsonLowerBound": 0.005215910778690411
}
},
"Cambricon_mlu-370-x4": {
"vllm": {
"framework": "vllm",
"modelCount": 25953,
"modelCount": 25961,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3186,
"successRate": 0.12276037452317651,
"wilsonLowerBound": 0.1188235595334533
"successCount": 3187,
"successRate": 0.12276106467393398,
"wilsonLowerBound": 0.11882483798483397
},
"vllm-customized": {
"framework": "vllm-customized",
@@ -115,23 +115,23 @@
},
"vllm-mlu": {
"framework": "vllm-mlu",
"modelCount": 9391,
"modelCount": 9407,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 4077,
"successRate": 0.434139069321691,
"wilsonLowerBound": 0.424143358656614
"successCount": 4084,
"successRate": 0.43414478579781013,
"wilsonLowerBound": 0.4241575354389301
}
},
"Cambricon_mlu-370-x8": {
"vllm": {
"framework": "vllm",
"modelCount": 24027,
"modelCount": 24030,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3770,
"successRate": 0.15690681316851876,
"wilsonLowerBound": 0.1523626843060331
"successRate": 0.15688722430295465,
"wilsonLowerBound": 0.1523436123834872
},
"vllm-customized": {
"framework": "vllm-customized",
@@ -144,50 +144,50 @@
},
"vllm-mlu": {
"framework": "vllm-mlu",
"modelCount": 39385,
"modelCount": 39427,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 9934,
"successRate": 0.25222800558588293,
"wilsonLowerBound": 0.2479631553725973
"successCount": 9945,
"successRate": 0.25223831384584167,
"wilsonLowerBound": 0.24797566376420394
}
},
"Iluvatar_bi-100": {
"transformers": {
"framework": "transformers",
"modelCount": 22914,
"modelCount": 22915,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3026,
"successRate": 0.1320590032294667,
"wilsonLowerBound": 0.1277369743815997
"successRate": 0.1320532402356535,
"wilsonLowerBound": 0.12773138638590842
},
"vllm": {
"framework": "vllm",
"modelCount": 86481,
"modelCount": 86485,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 9762,
"successRate": 0.11288028584313317,
"wilsonLowerBound": 0.11078836443747443
"successRate": 0.11287506504018038,
"wilsonLowerBound": 0.11078323440993057
},
"vllm-patch-tokenizer": {
"framework": "vllm-patch-tokenizer",
"modelCount": 38609,
"modelCount": 38614,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 4738,
"successRate": 0.12271750110077961,
"wilsonLowerBound": 0.11948206919093914
"successRate": 0.12270161081473041,
"wilsonLowerBound": 0.11946656975844437
},
"vllm_fix_tokenizer": {
"framework": "vllm_fix_tokenizer",
"modelCount": 26113,
"modelCount": 26133,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3210,
"successRate": 0.1229272776011948,
"wilsonLowerBound": 0.11900002209957782
"successCount": 3213,
"successRate": 0.12294799678567328,
"wilsonLowerBound": 0.11902193180839078
}
},
"Iluvatar_bi-150": {
@@ -202,30 +202,30 @@
},
"transformers": {
"framework": "transformers",
"modelCount": 5051,
"modelCount": 5054,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 562,
"successRate": 0.11126509602058998,
"wilsonLowerBound": 0.10288651808446177
"successRate": 0.111199050257222,
"wilsonLowerBound": 0.10282517065480763
},
"vllm": {
"framework": "vllm",
"modelCount": 112422,
"modelCount": 112425,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 16502,
"successRate": 0.14678621622102436,
"wilsonLowerBound": 0.1447295643963279
"successRate": 0.14678229931065154,
"wilsonLowerBound": 0.14472569775049507
},
"vllm_0_17_0_corex_4_4_0": {
"framework": "vllm_0_17_0_corex_4_4_0",
"modelCount": 15606,
"modelCount": 15613,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3187,
"successRate": 0.2042163270536973,
"wilsonLowerBound": 0.19796458864174885
"successCount": 3188,
"successRate": 0.2041888170114648,
"wilsonLowerBound": 0.19793878702309592
},
"vllm_fix_tokenizer": {
"framework": "vllm_fix_tokenizer",
@@ -238,12 +238,12 @@
},
"vllm_tokenizer_patch": {
"framework": "vllm_tokenizer_patch",
"modelCount": 6020,
"modelCount": 6025,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 469,
"successRate": 0.07790697674418605,
"wilsonLowerBound": 0.07140227111708827
"successRate": 0.07784232365145229,
"wilsonLowerBound": 0.07134281712711635
}
},
"Iluvatar_mrv-100": {
@@ -258,81 +258,81 @@
},
"vllm": {
"framework": "vllm",
"modelCount": 29851,
"modelCount": 29865,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 5085,
"successRate": 0.1703460520585575,
"wilsonLowerBound": 0.16612380778830965
"successCount": 5092,
"successRate": 0.17050058597019924,
"wilsonLowerBound": 0.16627776573653213
}
},
"Kunlunxin_p-800": {
"vllm": {
"framework": "vllm",
"modelCount": 60731,
"modelCount": 60733,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 12375,
"successRate": 0.20376743343597176,
"wilsonLowerBound": 0.2005826178036079
"successCount": 12376,
"successRate": 0.20377718867831326,
"wilsonLowerBound": 0.20059236750741152
},
"vllm_fix_tokenizer": {
"framework": "vllm_fix_tokenizer",
"modelCount": 5963,
"modelCount": 5978,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 1765,
"successRate": 0.29599195036055675,
"wilsonLowerBound": 0.2845397771768032
"successCount": 1768,
"successRate": 0.29575108732017397,
"wilsonLowerBound": 0.28431600169064586
},
"vllm_tokenizer_patch": {
"framework": "vllm_tokenizer_patch",
"modelCount": 11689,
"modelCount": 11706,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3070,
"successRate": 0.26264008897253827,
"wilsonLowerBound": 0.2547411183645426
"successCount": 3073,
"successRate": 0.26251494959849647,
"wilsonLowerBound": 0.2546229223317426
}
},
"MetaX_c-500": {
"vllm": {
"framework": "vllm",
"modelCount": 58107,
"modelCount": 58155,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 15820,
"successRate": 0.2722563546560655,
"wilsonLowerBound": 0.26865223630882573
"successCount": 15830,
"successRate": 0.2722035938440375,
"wilsonLowerBound": 0.26860117980036025
}
},
"Mthreads_s4000": {
"llamacpp": {
"framework": "llamacpp",
"modelCount": 45825,
"modelCount": 45831,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 19871,
"successRate": 0.4336279323513366,
"wilsonLowerBound": 0.42909620642371943
"successCount": 19877,
"successRate": 0.43370207937858657,
"wilsonLowerBound": 0.42917055264073056
},
"vllm": {
"framework": "vllm",
"modelCount": 34986,
"modelCount": 35052,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 2968,
"successRate": 0.08483393357342937,
"wilsonLowerBound": 0.08195958326032747
"successRate": 0.08467419833390391,
"wilsonLowerBound": 0.08180502284967742
},
"vllm_tokenizer_patch": {
"framework": "vllm_tokenizer_patch",
"modelCount": 203,
"modelCount": 211,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 0,
"successRate": 0.0,
"wilsonLowerBound": 1.702505035850099e-18
"wilsonLowerBound": 0.0
}
},
"Sunrise_pt-200-x1": {
@@ -356,12 +356,12 @@
},
"vllm_fix_tokenizer": {
"framework": "vllm_fix_tokenizer",
"modelCount": 5360,
"modelCount": 5367,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 2244,
"successRate": 0.41865671641791047,
"wilsonLowerBound": 0.40551212475467374
"successCount": 2248,
"successRate": 0.4188559716787777,
"wilsonLowerBound": 0.4057188914791776
}
},
"Vastai_va16": {
@@ -376,161 +376,161 @@
},
"vllm_fix_tokenizer": {
"framework": "vllm_fix_tokenizer",
"modelCount": 9834,
"modelCount": 9857,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 3381,
"successRate": 0.3438071995118975,
"wilsonLowerBound": 0.33448201885340845
"successCount": 3390,
"successRate": 0.3439180277975043,
"wilsonLowerBound": 0.3346028962130698
}
},
"hygon_k100-ai": {
"llamacpp": {
"framework": "llamacpp",
"modelCount": 25567,
"modelCount": 25585,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 12857,
"successRate": 0.5028747995462901,
"wilsonLowerBound": 0.49674597776775475
"successCount": 12865,
"successRate": 0.5028336916161814,
"wilsonLowerBound": 0.4967070292687963
},
"vllm": {
"framework": "vllm",
"modelCount": 60420,
"modelCount": 60431,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 5862,
"successRate": 0.09702085402184707,
"wilsonLowerBound": 0.0946862739354135
"successCount": 5864,
"successRate": 0.09703628932170574,
"wilsonLowerBound": 0.0947017509028485
},
"vllm-patch-tokenizer": {
"framework": "vllm-patch-tokenizer",
"modelCount": 11994,
"modelCount": 12006,
"officialConfigError": null,
"officialConfigValid": true,
"successCount": 1622,
"successRate": 0.13523428380857094,
"wilsonLowerBound": 0.1292307288142234
"successCount": 1628,
"successRate": 0.13559886723305015,
"wilsonLowerBound": 0.1295911942318147
}
}
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-10T15:22:18.248614+00:00",
"generatedAt": "2026-09-10T15:27:32.924657+00:00",
"gpuStats": {
"Ascend_910-b3": {
"available": true,
"backlogHours": 1077.1384615384616,
"backlogHours": 1077.0923076923077,
"canVerify": true,
"error": null,
"gpu": "Ascend_910-b3",
"healthFactor": 1.0,
"maxConcurrentTasks": 8,
"qualityFactor": 1.3023494599865717,
"queueFactor": 0.9508142077470051,
"queueWeight": 0.9508142077470051,
"recentSuccess": 51,
"recentSuccessRate": 0.3923076923076923,
"qualityFactor": 1.504632795098374,
"queueFactor": 0.9512896617649836,
"queueWeight": 0.9512896617649836,
"recentSuccess": 54,
"recentSuccessRate": 0.4153846153846154,
"recentTerminal": 130,
"recentWilsonLowerBound": 0.3126200051001405,
"running": 8,
"selectionWeight": 1.238292370006872,
"recentWilsonLowerBound": 0.3342905969843901,
"running": 7,
"selectionWeight": 1.431341622729634,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 21.666666666666668,
"waiting": 23338
"waiting": 23337
},
"Ascend_910-b4": {
"available": true,
"backlogHours": 973.3136094674555,
"backlogHours": 996.8363636363637,
"canVerify": true,
"error": null,
"gpu": "Ascend_910-b4",
"healthFactor": 1.0,
"maxConcurrentTasks": 8,
"qualityFactor": 0.9886513030818428,
"queueFactor": 0.965380391416169,
"queueWeight": 0.965380391416169,
"recentSuccess": 58,
"recentSuccessRate": 0.3431952662721893,
"recentTerminal": 169,
"recentWilsonLowerBound": 0.27581302397641744,
"running": 8,
"selectionWeight": 0.954424581943255,
"qualityFactor": 1.087693595098827,
"queueFactor": 0.9624033694620674,
"queueWeight": 0.9624033694620674,
"recentSuccess": 59,
"recentSuccessRate": 0.3575757575757576,
"recentTerminal": 165,
"recentWilsonLowerBound": 0.2884481904837394,
"running": 7,
"selectionWeight": 1.0467999808654207,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 28.166666666666668,
"waiting": 27415
"throughputPerHour": 27.5,
"waiting": 27413
},
"Biren_166m": {
"available": true,
"backlogHours": 128.98536585365855,
"backlogHours": 132.8140703517588,
"canVerify": true,
"error": null,
"gpu": "Biren_166m",
"healthFactor": 1.0,
"maxConcurrentTasks": 8,
"qualityFactor": 0.24647406477127937,
"qualityFactor": 0.24614996127901195,
"queueFactor": 1.3,
"queueWeight": 1.3,
"recentSuccess": 40,
"recentSuccessRate": 0.1951219512195122,
"recentTerminal": 205,
"recentWilsonLowerBound": 0.14668991695050224,
"recentSuccess": 39,
"recentSuccessRate": 0.19597989949748743,
"recentTerminal": 199,
"recentWilsonLowerBound": 0.14680692645920698,
"running": 8,
"selectionWeight": 0.32041628420266316,
"selectionWeight": 0.3199949496627155,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 34.166666666666664,
"waiting": 4407
"throughputPerHour": 33.166666666666664,
"waiting": 4405
},
"Cambricon_mlu-370-x4": {
"available": true,
"backlogHours": 1021.7349397590361,
"backlogHours": 1046.5185185185185,
"canVerify": true,
"error": null,
"gpu": "Cambricon_mlu-370-x4",
"healthFactor": 1.0,
"maxConcurrentTasks": 7,
"qualityFactor": 0.13799428842124495,
"queueFactor": 0.9583753964830888,
"queueWeight": 0.9583753964830888,
"recentSuccess": 15,
"recentSuccessRate": 0.18072289156626506,
"recentTerminal": 83,
"recentWilsonLowerBound": 0.1126926689223679,
"running": 7,
"selectionWeight": 0.13225033087811233,
"qualityFactor": 0.11984357842622166,
"queueFactor": 0.9554075700017707,
"queueWeight": 0.9554075700017707,
"recentSuccess": 14,
"recentSuccessRate": 0.1728395061728395,
"recentTerminal": 81,
"recentWilsonLowerBound": 0.10584307626279289,
"running": 6,
"selectionWeight": 0.11449946204451307,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 13.833333333333334,
"waiting": 14134
"throughputPerHour": 13.5,
"waiting": 14128
},
"Cambricon_mlu-370-x8": {
"available": true,
"backlogHours": 432.59016393442624,
"backlogHours": 418.85714285714283,
"canVerify": true,
"error": null,
"gpu": "Cambricon_mlu-370-x8",
"healthFactor": 1.0,
"maxConcurrentTasks": 7,
"qualityFactor": 0.3773376449381269,
"queueFactor": 1.0902469907359134,
"queueWeight": 1.0902469907359134,
"recentSuccess": 30,
"recentSuccessRate": 0.2459016393442623,
"recentTerminal": 122,
"recentWilsonLowerBound": 0.17802151523932291,
"qualityFactor": 0.4493221962021021,
"queueFactor": 1.0960763996917333,
"queueWeight": 1.0960763996917333,
"recentSuccess": 33,
"recentSuccessRate": 0.2619047619047619,
"recentTerminal": 126,
"recentWilsonLowerBound": 0.19299483361516778,
"running": 7,
"selectionWeight": 0.4113912318851694,
"selectionWeight": 0.4924914551147827,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 20.333333333333332,
"throughputPerHour": 21.0,
"waiting": 8796
},
"Iluvatar_bi-100": {
"available": true,
"backlogHours": 66.90140845070422,
"backlogHours": 67.85714285714286,
"canVerify": true,
"error": null,
"gpu": "Iluvatar_bi-100",
@@ -539,197 +539,197 @@
"qualityFactor": 0.05,
"queueFactor": 1.3,
"queueWeight": 1.3,
"recentSuccess": 17,
"recentSuccessRate": 0.07981220657276995,
"recentTerminal": 213,
"recentWilsonLowerBound": 0.05042523581880395,
"recentSuccess": 16,
"recentSuccessRate": 0.0761904761904762,
"recentTerminal": 210,
"recentWilsonLowerBound": 0.047438977513945844,
"running": 1,
"selectionWeight": 0.065,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 35.5,
"submissionEligible": false,
"throughputPerHour": 35.0,
"waiting": 2375
},
"Iluvatar_bi-150": {
"available": true,
"backlogHours": 117.6923076923077,
"backlogHours": 120.15706806282724,
"canVerify": true,
"error": null,
"gpu": "Iluvatar_bi-150",
"healthFactor": 1.0,
"maxConcurrentTasks": 50,
"qualityFactor": 0.0647465175359186,
"qualityFactor": 0.0676123502754561,
"queueFactor": 1.3,
"queueWeight": 1.3,
"recentSuccess": 23,
"recentSuccessRate": 0.11794871794871795,
"recentTerminal": 195,
"recentWilsonLowerBound": 0.07989354857234383,
"recentSuccessRate": 0.12041884816753927,
"recentTerminal": 191,
"recentWilsonLowerBound": 0.08159575665748214,
"running": 1,
"selectionWeight": 0.08417047279669419,
"selectionWeight": 0.08789605535809294,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 32.5,
"throughputPerHour": 31.833333333333332,
"waiting": 3825
},
"Iluvatar_mrv-100": {
"available": true,
"backlogHours": 1855.4399999999998,
"backlogHours": 1932.375,
"canVerify": true,
"error": null,
"gpu": "Iluvatar_mrv-100",
"healthFactor": 1.0,
"maxConcurrentTasks": 2,
"qualityFactor": 0.18117436275647994,
"queueFactor": 0.8763333818567766,
"queueWeight": 0.8763333818567766,
"recentSuccess": 11,
"recentSuccessRate": 0.22,
"recentTerminal": 50,
"recentWilsonLowerBound": 0.1275378622430229,
"qualityFactor": 0.25508215803483625,
"queueFactor": 0.8714390238547999,
"queueWeight": 0.8714390238547999,
"recentSuccess": 12,
"recentSuccessRate": 0.25,
"recentTerminal": 48,
"recentWilsonLowerBound": 0.1492048880971003,
"running": 2,
"selectionWeight": 0.15876914202013248,
"selectionWeight": 0.2222885468006535,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 8.333333333333334,
"waiting": 15462
"throughputPerHour": 8.0,
"waiting": 15459
},
"Kunlunxin_p-800": {
"available": true,
"backlogHours": 692.1063829787234,
"backlogHours": 684.8210526315789,
"canVerify": true,
"error": null,
"gpu": "Kunlunxin_p-800",
"healthFactor": 1.0,
"maxConcurrentTasks": 7,
"qualityFactor": 0.4210453294615219,
"queueFactor": 1.0160392014722135,
"queueWeight": 1.0160392014722135,
"recentSuccess": 25,
"recentSuccessRate": 0.26595744680851063,
"recentTerminal": 94,
"recentWilsonLowerBound": 0.1871148582279006,
"qualityFactor": 0.45523063755852883,
"queueFactor": 1.0181555905320407,
"queueWeight": 1.0181555905320407,
"recentSuccess": 26,
"recentSuccessRate": 0.2736842105263158,
"recentTerminal": 95,
"recentWilsonLowerBound": 0.19414427878715157,
"running": 7,
"selectionWeight": 0.42779856032968977,
"selectionWeight": 0.4634956186116813,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 15.666666666666666,
"throughputPerHour": 15.833333333333334,
"waiting": 10843
},
"MetaX_c-500": {
"available": true,
"backlogHours": 599.7894736842105,
"backlogHours": 584.3076923076923,
"canVerify": true,
"error": null,
"gpu": "MetaX_c-500",
"healthFactor": 1.0,
"maxConcurrentTasks": 4,
"qualityFactor": 0.3668089035745469,
"queueFactor": 1.0380937271106694,
"queueWeight": 1.0380937271106694,
"recentSuccess": 28,
"recentSuccessRate": 0.24561403508771928,
"recentTerminal": 114,
"recentWilsonLowerBound": 0.17574622644308974,
"qualityFactor": 0.3783077320646403,
"queueFactor": 1.0426882394039771,
"queueWeight": 1.0426882394039771,
"recentSuccess": 29,
"recentSuccessRate": 0.24786324786324787,
"recentTerminal": 117,
"recentWilsonLowerBound": 0.1784782855379003,
"running": 4,
"selectionWeight": 0.3807820218490795,
"selectionWeight": 0.3944570230993913,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 19.0,
"waiting": 11396
"throughputPerHour": 19.5,
"waiting": 11394
},
"Mthreads_s4000": {
"available": true,
"backlogHours": 1907.2835820895523,
"backlogHours": 1879.0588235294117,
"canVerify": true,
"error": null,
"gpu": "Mthreads_s4000",
"healthFactor": 1.0,
"maxConcurrentTasks": 4,
"qualityFactor": 1.025898745090207,
"queueFactor": 0.8727183388955203,
"queueWeight": 0.8727183388955203,
"qualityFactor": 0.9878212979600234,
"queueFactor": 0.8751039805539271,
"queueWeight": 0.8751039805539271,
"recentSuccess": 26,
"recentSuccessRate": 0.3880597014925373,
"recentTerminal": 67,
"recentWilsonLowerBound": 0.2804887105692933,
"recentSuccessRate": 0.38235294117647056,
"recentTerminal": 68,
"recentWilsonLowerBound": 0.2760927529710457,
"running": 4,
"selectionWeight": 0.8953206486901243,
"selectionWeight": 0.8644463499207634,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 11.166666666666666,
"waiting": 21298
"throughputPerHour": 11.333333333333334,
"waiting": 21296
},
"Sunrise_pt-200-x1": {
"available": true,
"backlogHours": 3852.2727272727275,
"backlogHours": 3684.782608695652,
"canVerify": true,
"error": null,
"gpu": "Sunrise_pt-200-x1",
"healthFactor": 1.0,
"maxConcurrentTasks": 2,
"qualityFactor": 1.2533531071700206,
"queueFactor": 0.7853781937174696,
"queueWeight": 0.7853781937174696,
"qualityFactor": 1.120483029902168,
"queueFactor": 0.7910226788918543,
"queueWeight": 0.7910226788918543,
"recentSuccess": 11,
"recentSuccessRate": 0.5,
"recentTerminal": 22,
"recentWilsonLowerBound": 0.3072180469247956,
"recentSuccessRate": 0.4782608695652174,
"recentTerminal": 23,
"recentWilsonLowerBound": 0.29236869539945726,
"running": 2,
"selectionWeight": 0.9843561993993689,
"selectionWeight": 0.8863274879660746,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 3.6666666666666665,
"throughputPerHour": 3.8333333333333335,
"waiting": 14125
},
"Vastai_va16": {
"available": true,
"backlogHours": 847.0140845070422,
"backlogHours": 859.3714285714286,
"canVerify": true,
"error": null,
"gpu": "Vastai_va16",
"healthFactor": 1.0,
"maxConcurrentTasks": 16,
"qualityFactor": 0.6414138264514839,
"queueFactor": 0.9857182512242489,
"queueWeight": 0.9857182512242489,
"qualityFactor": 0.660853348951212,
"queueFactor": 0.9840645319666167,
"queueWeight": 0.9840645319666167,
"recentSuccess": 23,
"recentSuccessRate": 0.323943661971831,
"recentTerminal": 71,
"recentWilsonLowerBound": 0.22657056384385282,
"recentSuccessRate": 0.32857142857142857,
"recentTerminal": 70,
"recentWilsonLowerBound": 0.2299871182671781,
"running": 16,
"selectionWeight": 0.6322533153208106,
"selectionWeight": 0.6503223415342456,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 11.833333333333334,
"waiting": 10023
"throughputPerHour": 11.666666666666666,
"waiting": 10026
},
"hygon_k100-ai": {
"available": true,
"backlogHours": 485.3831775700935,
"backlogHours": 476.31192660550454,
"canVerify": true,
"error": null,
"gpu": "hygon_k100-ai",
"healthFactor": 1.0,
"maxConcurrentTasks": 6,
"qualityFactor": 2.5,
"queueFactor": 1.0715777430310858,
"queueWeight": 1.0715777430310858,
"recentSuccess": 55,
"recentSuccessRate": 0.514018691588785,
"recentTerminal": 107,
"recentWilsonLowerBound": 0.42048422635928223,
"running": 6,
"selectionWeight": 2.6789443575777145,
"queueFactor": 1.0751448952718026,
"queueWeight": 1.0751448952718026,
"recentSuccess": 56,
"recentSuccessRate": 0.5137614678899083,
"recentTerminal": 109,
"recentWilsonLowerBound": 0.4210714006745304,
"running": 5,
"selectionWeight": 2.6878622381795063,
"stale": false,
"submissionEligible": true,
"throughputPerHour": 17.833333333333332,
"waiting": 8656
"throughputPerHour": 18.166666666666668,
"waiting": 8653
}
},
"queueAttemptedAt": "2026-09-10T15:16:33.850208+00:00",
"queueAttemptedAt": "2026-09-10T15:27:32.924657+00:00",
"queueError": null,
"queueUpdatedAt": "2026-09-10T15:16:33.850208+00:00",
"queueUpdatedAt": "2026-09-10T15:27:32.924657+00:00",
"supportedGpus": [
"Vastai_va16",
"Iluvatar_mrv-100",

File diff suppressed because it is too large Load Diff

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-10T15:19:40.193833+00:00",
"lastSyncTime": "2026-09-10T15:19:40.009069+00:00",
"generatedAt": "2026-09-10T15:27:10.017950+00:00",
"lastSyncTime": "2026-09-10T15:27:09.820580+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -1596,11 +1596,11 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 7,
"failureBreakdown": {
"ambiguous_runtime": 18,
"ambiguous_runtime": 19,
"framework_architecture_unsupported": 7,
"参数/模板问题": 3
},
"failureCount": 28,
"failureCount": 29,
"failureRate": 1.0,
"framework": "vllm-mlu",
"pendingCount": 0,
@@ -1610,8 +1610,8 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 28,
"unresolvedFailureCount": 21
"total": 29,
"unresolvedFailureCount": 22
},
"Cambricon_mlu-370-x8|vllm|text-generation": {
"attributableFailureCount": 3,
@@ -1906,9 +1906,9 @@
},
"MetaX_c-500|vllm|text-generation": {
"attributableFailureCount": 45,
"decisionFailureRate": 0.9783,
"decisionSuccessRate": 0.0217,
"decisionTotal": 46,
"decisionFailureRate": 0.9574,
"decisionSuccessRate": 0.0426,
"decisionTotal": 47,
"failureBreakdown": {
"ambiguous_runtime": 7,
"backend_operator": 26,
@@ -1918,16 +1918,16 @@
"参数/模板问题": 9
},
"failureCount": 61,
"failureRate": 0.9839,
"failureRate": 0.9683,
"framework": "vllm",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 1,
"successRate": 0.0161,
"successCount": 2,
"successRate": 0.0317,
"targetGpu": "MetaX_c-500",
"taskType": "text-generation",
"total": 62,
"total": 63,
"unresolvedFailureCount": 16
},
"Mthreads_s4000|llamacpp|text-generation": {
@@ -2317,9 +2317,9 @@
},
"vllm": {
"attributableFailureCount": 239,
"decisionFailureRate": 0.9958,
"decisionSuccessRate": 0.0042,
"decisionTotal": 240,
"decisionFailureRate": 0.9917,
"decisionSuccessRate": 0.0083,
"decisionTotal": 241,
"failureBreakdown": {
"ambiguous_runtime": 172,
"backend_operator": 32,
@@ -2333,13 +2333,13 @@
"参数/模板问题": 24
},
"failureCount": 436,
"failureRate": 0.9977,
"failureRate": 0.9954,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 1,
"successCount": 1,
"successRate": 0.0023,
"total": 437,
"successCount": 2,
"successRate": 0.0046,
"total": 438,
"unresolvedFailureCount": 196
},
"vllm-mlu": {
@@ -2348,19 +2348,19 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 7,
"failureBreakdown": {
"ambiguous_runtime": 18,
"ambiguous_runtime": 19,
"framework_architecture_unsupported": 7,
"参数/模板问题": 3
},
"failureCount": 28,
"failureCount": 29,
"failureRate": 1.0,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 0,
"successRate": 0.0,
"total": 28,
"unresolvedFailureCount": 21
"total": 29,
"unresolvedFailureCount": 22
},
"vllm-patch-tokenizer": {
"attributableFailureCount": 3,
@@ -2443,7 +2443,7 @@
"unresolvedFailureCount": 1
}
},
"generatedAt": "2026-09-10T15:19:40.187828+00:00",
"generatedAt": "2026-09-10T15:27:10.014184+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 26,
@@ -2545,22 +2545,22 @@
"decisionSuccessRate": 0.1053,
"decisionTotal": 19,
"failureBreakdown": {
"ambiguous_runtime": 37,
"ambiguous_runtime": 38,
"framework_architecture_unsupported": 15,
"memory_capacity": 1,
"tokenizer_compatibility": 1,
"参数/模板问题": 18,
"验证失败": 22
},
"failureCount": 94,
"failureRate": 0.9792,
"failureCount": 95,
"failureRate": 0.9794,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 2,
"successRate": 0.0208,
"total": 96,
"unresolvedFailureCount": 77
"successRate": 0.0206,
"total": 97,
"unresolvedFailureCount": 78
},
"Iluvatar_bi-100": {
"attributableFailureCount": 0,
@@ -2651,9 +2651,9 @@
},
"MetaX_c-500": {
"attributableFailureCount": 45,
"decisionFailureRate": 0.8182,
"decisionSuccessRate": 0.1818,
"decisionTotal": 55,
"decisionFailureRate": 0.8036,
"decisionSuccessRate": 0.1964,
"decisionTotal": 56,
"failureBreakdown": {
"ambiguous_runtime": 7,
"backend_operator": 26,
@@ -2664,13 +2664,13 @@
"验证失败": 48
},
"failureCount": 117,
"failureRate": 0.9213,
"failureRate": 0.9141,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 10,
"successRate": 0.0787,
"total": 127,
"successCount": 11,
"successRate": 0.0859,
"total": 128,
"unresolvedFailureCount": 72
},
"Mthreads_s4000": {
@@ -2931,9 +2931,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 2
"ambiguous_runtime": 3
},
"failureCount": 2,
"failureCount": 3,
"failureRate": 1.0,
"framework": "vllm-mlu",
"modelType": "phi3",
@@ -2945,8 +2945,8 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 2,
"unresolvedFailureCount": 2
"total": 3,
"unresolvedFailureCount": 3
},
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|qwen2|compressed-tensors": {
"attributableFailureCount": 0,
@@ -3730,25 +3730,25 @@
},
"MetaX_c-500|vllm|text-generation|starcoder2|compressed-tensors": {
"attributableFailureCount": 7,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 7,
"decisionFailureRate": 0.875,
"decisionSuccessRate": 0.125,
"decisionTotal": 8,
"failureBreakdown": {
"backend_operator": 7
},
"failureCount": 7,
"failureRate": 1.0,
"failureRate": 0.875,
"framework": "vllm",
"modelType": "starcoder2",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"successCount": 1,
"successRate": 0.125,
"targetGpu": "MetaX_c-500",
"taskType": "text-generation",
"total": 7,
"total": 8,
"unresolvedFailureCount": 0
},
"Mthreads_s4000|vllm|text-generation|glm_ocr|fp8": {
@@ -4147,11 +4147,11 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 4,
"failureBreakdown": {
"ambiguous_runtime": 14,
"ambiguous_runtime": 15,
"framework_architecture_unsupported": 4,
"参数/模板问题": 1
},
"failureCount": 19,
"failureCount": 20,
"failureRate": 1.0,
"framework": "vllm-mlu",
"lastPlatformFailureAt": null,
@@ -4163,8 +4163,8 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 19,
"unresolvedFailureCount": 15
"total": 20,
"unresolvedFailureCount": 16
},
"Cambricon_mlu-370-x8|vllm|text-generation": {
"attributableFailureCount": 3,
@@ -4713,9 +4713,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
"ambiguous_runtime": 2
},
"failureCount": 1,
"failureCount": 2,
"failureRate": 1.0,
"framework": "vllm-mlu",
"lastTerminalAt": "2026-09-08T20:25:19.301781+00:00",
@@ -4728,8 +4728,8 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
"total": 2,
"unresolvedFailureCount": 2
},
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|qwen2|compressed-tensors": {
"attributableFailureCount": 0,
@@ -5404,9 +5404,9 @@
"decisionSuccessRate": 0.0,
"decisionTotal": 0,
"failureBreakdown": {
"ambiguous_runtime": 1
"ambiguous_runtime": 2
},
"failureCount": 1,
"failureCount": 2,
"failureRate": 1.0,
"framework": "vllm-mlu",
"loadSizeLog2Bucket": 33,
@@ -5419,8 +5419,8 @@
"successRate": 0.0,
"targetGpu": "Cambricon_mlu-370-x8",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 1
"total": 2,
"unresolvedFailureCount": 2
},
"Cambricon_mlu-370-x8|vllm-mlu|text-generation|qwen2|compressed-tensors|30": {
"attributableFailureCount": 0,
@@ -6810,14 +6810,14 @@
},
"MetaX_c-500|vllm|text-generation|starcoder2|compressed-tensors|34": {
"attributableFailureCount": 2,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 2,
"decisionFailureRate": 0.6667,
"decisionSuccessRate": 0.3333,
"decisionTotal": 3,
"failureBreakdown": {
"backend_operator": 2
},
"failureCount": 2,
"failureRate": 1.0,
"failureRate": 0.6667,
"framework": "vllm",
"loadSizeLog2Bucket": 34,
"modelType": "starcoder2",
@@ -6825,11 +6825,11 @@
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"successCount": 1,
"successRate": 0.3333,
"targetGpu": "MetaX_c-500",
"taskType": "text-generation",
"total": 2,
"total": 3,
"unresolvedFailureCount": 0
},
"Mthreads_s4000|vllm|text-generation|glm_ocr|fp8|30": {
@@ -7121,15 +7121,15 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 1598,
"totalRecords": 1685,
"terminalRecords": 1600,
"totalRecords": 1687,
"totals": {
"attributableFailureCount": 419,
"decisionFailureRate": 0.8915,
"decisionSuccessRate": 0.1085,
"decisionTotal": 470,
"decisionFailureRate": 0.8896,
"decisionSuccessRate": 0.1104,
"decisionTotal": 471,
"failureBreakdown": {
"ambiguous_runtime": 354,
"ambiguous_runtime": 355,
"backend_operator": 39,
"framework_architecture_unsupported": 272,
"memory_capacity": 11,
@@ -7141,49 +7141,49 @@
"参数/模板问题": 98,
"验证失败": 673
},
"failureCount": 1547,
"failureRate": 0.9681,
"failureCount": 1548,
"failureRate": 0.9675,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 3,
"successCount": 51,
"successRate": 0.0319,
"total": 1598,
"unresolvedFailureCount": 1125
"successCount": 52,
"successRate": 0.0325,
"total": 1600,
"unresolvedFailureCount": 1126
},
"warnings": [
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Sunrise_pt-200-x1 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Iluvatar_mrv-100|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x4|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Vastai_va16|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Vastai_va16|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b3|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|vllm-mlu|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm-patch-tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Mthreads_s4000|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x4|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
"组合 Cambricon_mlu-370-x8|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 MetaX_c-500|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Vastai_va16|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Sunrise_pt-200-x1|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Ascend_910-b4|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Iluvatar_bi-150|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Vastai_va16|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 hygon_k100-ai|vllm-patch-tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 1685,
"summarizedRecords": 1687,
"version": 1
}

View File

@@ -116,13 +116,13 @@
"unsloth/LFM2-2.6B-Exp-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/unsloth/LFM2-2.6B-Exp-GGUF/resolve/master/config.json (status=404)",
"z-lab/Qwen3.8-27B-DFlash2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/z-lab/Qwen3.8-27B-DFlash2-GGUF/resolve/master/config.json (status=404)"
},
"architectureModelConfigsComplete": 125,
"architectureModelConfigsComplete": 123,
"architectureOnly": false,
"architecturePolicySkipped": {
"frameworkCatalogUnknown": 2,
"frameworkContextUnknown": 382,
"modelArchitectureUnknown": 166,
"noMatchingBlock": 812,
"frameworkContextUnknown": 377,
"modelArchitectureUnknown": 168,
"noMatchingBlock": 810,
"partiallyBlockedFrameworkSet": 102,
"runningMatchedProtected": 0,
"submissionContextMismatch": 0,
@@ -165,12 +165,12 @@
"repositorySizeErrors": {
"empero-ai/Qwen3.8-9B-GGUF": "recursive_repository_size_incomplete"
},
"repositorySizesComplete": 266,
"repositorySizesComplete": 267,
"skipped": {
"fitsKnownCapacity": 1080,
"gpuCapacityUnknown": 0,
"repositorySizeUnknown": 2
},
"stopErrors": [],
"uniqueModels": 267
"uniqueModels": 268
}

View File

@@ -280,6 +280,7 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T15:41:33.598990+00:00", "modelId": "RedHatAI/starcoder2-3b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4093180600, "estimatedRequiredGiB": 4.578, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 4096479058, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 3181366272, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4096479058}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T06:46:06.888683+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712801", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T20:25:19.301781+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020365960, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022810352, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022810352}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T06:46:06.887685+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712810", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T18:40:55.416614+00:00", "modelId": "RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3SmallForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8630134648, "estimatedRequiredGiB": 9.647, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3small", "modelscopeFileSize": 8632060880, "modelscopeLicense": "mit", "modelscopeParams": 7803314176, "modelscopeTags": ["license:mit", "model_type:phi3small", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 8632060880}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T06:46:06.886454+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712804", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-10T15:27:09.820580+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14293335560, "estimatedRequiredGiB": 15.977, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 14295942065, "modelscopeLicense": "mit", "modelscopeParams": 13960238080, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 14295942065}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T06:46:06.810137+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712809", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-10T13:14:20.004644+00:00", "modelId": "RedHatAI/starcoder2-3b-FP8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3484485728, "estimatedRequiredGiB": 3.898, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 3487786133, "modelscopeLicense": "other", "modelscopeParams": 3181366272, "modelscopeTags": ["license:other", "model_type:starcoder2", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 3487786133}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T06:46:06.808048+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712807", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T20:58:55.412656+00:00", "modelId": "RedHatAI/Qwen2.5-1.5B-quantized.w4a16", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1611497856, "estimatedRequiredGiB": 1.819, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 1627385309, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1781159424, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:chat", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 1627385309}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T06:46:06.804740+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712806", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T15:41:33.599000+00:00", "modelId": "neuralmagic/Llama-3.2-1B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2023929984, "estimatedRequiredGiB": 2.272, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 2033083373, "modelscopeLicense": "llama3.2", "modelscopeParams": 1498482688, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2033083377}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T06:46:06.795745+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712800", "taskType": "text-generation", "verifyResult": -1}
@@ -297,4 +298,3 @@
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-08T03:17:11.503517+00:00", "modelId": "ysqlian/YZH_Model", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T03:11:21+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4487542", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "model_load", "failureAction": "inspect_root_exception_then_llm_if_unknown", "failureCategory": "model_load", "failureClassificationReason": "broad_load_error", "failureCode": "MODEL_LOAD_FAILED", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-08T03:17:11.503505+00:00", "modelId": "nightmedia/Qwen3-4B-Element8-Eva-Heretic-mxfp4-mlx", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T03:05:21+00:00", "targetGpu": "Vastai_va16", "taskId": "4079944", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "", "lastSyncTime": "2026-09-08T03:04:50.903963+00:00", "modelId": "mlx-community/Qwen3.5-35B-A3B-OptiQ-4bit-REAP-19B", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T03:01:23+00:00", "targetGpu": "Biren_166m", "taskId": "4610405", "taskType": "text-generation", "verifyResult": -1}
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "framework": "vllm", "lastSyncTime": "2026-09-08T03:04:50.903987+00:00", "modelId": "microsoft/Phi-3-small-128k-instruct", "modelProfile": {}, "outcome": "failed", "status": "success", "submitTime": "2026-09-08T03:01:22+00:00", "targetGpu": "Vastai_va16", "taskId": "4079318", "taskType": "text-generation", "verifyResult": -1}

View File

@@ -645,6 +645,7 @@
{"batchId": "907987834148411ea675a5e6a75ca42a", "completedAt": "2026-09-10T14:46:38.200736+00:00", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-10T14:46:30.271611+00:00", "framework": "vllm", "intentId": "28a76d8eaec247d6bb2628db7901d520", "lastModified": "2026-09-10T05:57:11+00:00", "modelAddress": "https://modelscope.cn/models/CohereLabs/tiny-aya-en-thinker", "reason": null, "reconciledAt": "2026-09-10T15:27:09.087043+00:00", "repoId": "CohereLabs/tiny-aya-en-thinker", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4761530", "taskType": "text-generation"}
{"batchId": "907987834148411ea675a5e6a75ca42a", "completedAt": "2026-09-10T14:46:38.200738+00:00", "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configSource": "modelhub_live", "createdAt": "2026-09-10T14:46:30.271653+00:00", "framework": "vllm", "intentId": "0a351039ed9e443999fbb111a9ba1df7", "lastModified": "2026-08-11T14:42:47+00:00", "modelAddress": "https://modelscope.cn/models/BAAI/AREX-Turbo", "reason": null, "reconciledAt": "2026-09-10T15:27:09.087175+00:00", "repoId": "BAAI/AREX-Turbo", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4761531", "taskType": "text-generation"}
{"batchId": "c3b3a3e3df6e4a6c8f33bece0f8ec90d", "completedAt": "2026-09-10T15:13:34.008724+00:00", "configFingerprint": "e4ab8f6cadfae4845d02aa9c1f89de5907e1f311ae817b1e8e370b501bd7f2e9", "configSource": "modelhub_live", "createdAt": "2026-09-10T15:13:25.575035+00:00", "framework": "llamacpp", "intentId": "4831e017dc354a2da3332acef4a9cf4d", "lastModified": "2026-09-10T14:58:25+00:00", "modelAddress": "https://modelscope.cn/models/webAI-Official/TwIL-LM3", "reason": null, "reconciledAt": "2026-09-10T15:27:09.087463+00:00", "repoId": "webAI-Official/TwIL-LM3", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "hygon_k100-ai", "taskId": "4761809", "taskType": "text-generation"}
{"batchId": "eb410d3235d340868c8cf6ab7434f833", "configFingerprint": "2be5f709184000d7e6bae759e29c6586ad3430b235ef3eb2dab86a5c2487a2f7", "configSource": "modelhub_live", "createdAt": "2026-09-10T15:29:42.957367+00:00", "framework": "llamacpp", "intentId": "b88d7ddd69e94be882df8a64076be3e2", "lastModified": "2026-09-10T14:58:25+00:00", "modelAddress": "https://modelscope.cn/models/webAI-Official/TwIL-LM3", "repoId": "webAI-Official/TwIL-LM3", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
{"batchId": "93abbd86b4564ff2b94a94e46d15997c", "completedAt": "2026-09-10T02:01:15.152981+00:00", "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configSource": "modelhub_live", "createdAt": "2026-09-10T02:01:07.997982+00:00", "framework": "vllm", "intentId": "f62de0b86d234667845793d4dabba8be", "repoId": "nv-community/Qwen3.8-27B-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "Iluvatar_mrv-100", "taskId": null, "taskType": "text-generation"}
{"batchId": "60b640169cee4eb6aeaf36aa752d68d4", "completedAt": "2026-09-09T21:08:33.584377+00:00", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-09T21:08:21.304731+00:00", "framework": "vllm_fix_tokenizer", "intentId": "31500eb429ef4b3cbc844e4096a6b5f6", "repoId": "nv-community/Qwen3.8-27B-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "Kunlunxin_p-800", "taskId": null, "taskType": "text-generation"}
{"batchId": "395e59f0f1df48fcb2e0269b69d9bce5", "completedAt": "2026-09-09T16:24:21.684805+00:00", "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configSource": "modelhub_live", "createdAt": "2026-09-09T16:24:14.160686+00:00", "framework": "vllm", "intentId": "e2523fa4e63645688fd5e7c0c7c1a5d5", "repoId": "nv-community/Qwen3.8-27B-NVFP4", "safeConfigVector": {"dtype": "-", "gpuMemoryUtilization": 0.9, "gpuNum": 1, "maxModelLen": 4096, "tensorParallel": 1}, "status": "uniqueness_rejected", "targetGpu": "Biren_166m", "taskId": null, "taskType": "text-generation"}

View File

@@ -1,24 +1,24 @@
{
"agentVersion": "2026.09.04.4",
"checksums": {
".modelhub_state/architecture_compatibility_blacklist.json": "235d7475bd098befbcf8c18f39eb0741f5cfa3ce9428ecc170e36cb43290f04c",
".modelhub_state/architecture_compatibility_blacklist.json": "568ca073a7f43a3a3fb770afcc579db6ce7a5d45bbd87d171034761f422548c1",
".modelhub_state/architecture_history_backfill.json": "212e68ed32b8c9406e433da8bd736d14b608fd4dc8bce490120b00c43f527f40",
".modelhub_state/market_intelligence.json": "53ff670dc86f3341335610c4ce38b4717d729b906dd9a64eb0131b4db5933f43",
".modelhub_state/official_capabilities.json": "ee1e65d13607aca2ff9276b35341953e2234261c1645bbc5288009da97504605",
".modelhub_state/outcome_checkpoint.json": "108177bc497814844051121e936e96f0700fc5281361e3d9a5770d59f3746280",
".modelhub_state/queue_cleanup_latest.json": "a596175137b8adfa14ee62a71c9a7bc915b25b1d8f0f57451a25c2e257396039",
".modelhub_state/recent_outcomes.jsonl": "1ec5cb7fbc28c985edb48f86fba95f976f7d0a1ce5ff1f48a9023fbdc22e6e77",
".modelhub_state/market_intelligence.json": "e2624a55c80c0e1b5ed8e13203bd981a46a6a0ff2f1f701577a516c9f535aac7",
".modelhub_state/official_capabilities.json": "3d6016ce14b275ff12bdb7afe5652adbfdfb2b73a5248b463e0d284da831145b",
".modelhub_state/outcome_checkpoint.json": "58029266bb92b932d377336c9fccb3b6ebe1a82cdd97af970a66ed989f9c8b82",
".modelhub_state/queue_cleanup_latest.json": "c4c0c2f745124fbabf68b2aa082c5aeee0f45dbb41d40369e8064e674a33adbe",
".modelhub_state/recent_outcomes.jsonl": "9e9cb1dd0f2d6f58384a1234e9beec7ab6c7ae3a92d636f21215b319e8b90f3a",
".modelhub_state/recovery_active_tasks.jsonl": "758f4706189a2f3afe240425077b9475217166dd358888f2f3dac2eea1c46ca8",
".modelhub_state/recovery_intents.jsonl": "ff4fadd012130a96172439815e060600dba3b279610bbd31646053ea5dc7cca2",
".modelhub_state/recovery_intents.jsonl": "38d6d643700900c888bf36771e71c56c7952a68197c205c3b4beb73808f2083c",
".modelhub_state/routing_intelligence.json": "8901830c5979661bd91c7ff07bb088fddcd15c778651eb726d0d1d62df70e228",
".modelhub_state/submission_exclusions.jsonl": "b2de1125bf1f3e0597cd50cc30b563a5bfee39459e843767163eb72101cebc03",
".modelhub_state/worker_crashes.jsonl": "7fa483580e87374493226e80bcf40bf8044847acba70045c96fa175966a1a8ba",
"ledger/submissions.jsonl": "5cd9fb7f33b669b6f7abe4d47d9552e490231dcb1e0df0349f76c8eaac67656a",
"outcomes/submissions.jsonl": "5562d3c4413dc849f39649dd11745b00e8dbafdf04c96bc271a89549f9ef0c92"
"outcomes/submissions.jsonl": "c87ab0ad9e02c3f12d5e8f14b56e812e8fdfe26a8e271df3c60a3bfec5046918"
},
"generation": 5229,
"phase": "startup",
"generation": 5230,
"phase": "intent",
"schemaVersion": 1,
"updatedAt": "2026-09-10T15:27:09.162273+00:00",
"updatedAt": "2026-09-10T15:29:43.016437+00:00",
"writerId": "18252d8a8ef94333abd55823b2d81c42"
}

View File

@@ -173,7 +173,6 @@
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T11:53:42.913364+00:00", "modelId": "ornith-ai/Ornith-1.5-9B-MLX-8bit", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9514532079, "estimatedRequiredGiB": 10.658, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9536343971, "modelscopeLicense": null, "modelscopeParams": 2519020032, "modelscopeTags": ["model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9536343971}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T03:52:18.510578+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4660434", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T12:21:22.998937+00:00", "modelId": "ornith-ai/Ornith-1.5-9B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8818103892, "estimatedRequiredGiB": 9.883, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8842796317, "modelscopeLicense": "mit", "modelscopeParams": 6728625904, "modelscopeTags": ["license:mit", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 8842796317}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T04:17:45.285806+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4660827", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T12:39:38.309054+00:00", "modelId": "groxaxo/Qwen3.8-27B-GPTQ-Pro-4bit-g64-calib128", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 20034587064, "estimatedRequiredGiB": 22.413, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 20054851725, "modelscopeLicense": "apache-2.0", "modelscopeParams": 27781427952, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:qwen3.8", "custom_tag:qwen3.5-architecture", "custom_tag:gptq-pro", "custom_tag:gptq", "custom_tag:4-bit", "custom_tag:4bit", "custom_tag:mtp", "custom_tag:speculative-decoding", "custom_tag:text-generation", "custom_tag:conversational"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "gptq", "repositoryOnDiskBytes": 20054851725}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T04:38:16.005175+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4661079", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T13:08:06.614699+00:00", "modelId": "RedHatAI/starcoder2-15b-quantized.w8a8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785240496, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788612468, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "recentProfileFeedback": {"attributableFailureCount": 2, "consecutiveFailures": 2, "decisionFailureRate": 1.0, "decisionSuccessRate": 0.0, "decisionTotal": 2, "failureBreakdown": {"backend_operator": 2}, "failureCount": 2, "failureRate": 1.0, "framework": "vllm", "lastTerminalAt": "2026-09-05T10:14:38.509276+00:00", "modelType": "starcoder2", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "compressed-tensors", "successCount": 0, "successRate": 0.0, "targetGpu": "MetaX_c-500", "taskType": "text-generation", "total": 2, "unresolvedFailureCount": 0}, "repositoryOnDiskBytes": 17788612468}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T05:06:36.456528+00:00", "targetGpu": "MetaX_c-500", "taskId": "4661426", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-06T13:08:06.614723+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-FP8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14289052464, "estimatedRequiredGiB": 15.972, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 14291553638, "modelscopeLicense": "mit", "modelscopeParams": 13960238080, "modelscopeTags": ["license:mit", "model_type:phi3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 14291553638}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T05:06:36.493823+00:00", "targetGpu": "MetaX_c-500", "taskId": "4661437", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-06T14:42:01.034713+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T06:30:19.979390+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4662535", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-06T15:03:37.399221+00:00", "modelId": "cyankiwi/Apodex-1.1-mini-AWQ-INT4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 26119812080, "estimatedRequiredGiB": 29.233, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 26156887523, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 26156887523}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-06T07:02:35.712857+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4662918", "taskType": "text-generation", "verifyResult": null}
@@ -244,7 +243,6 @@
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-08T13:16:14.206207+00:00", "modelId": "OpenBMB/MiniCPM5-2B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1557c77dfe4d3dc1a68ae6b8349e7d094da7eec7752d29442a9d6d27e5f13121", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557096, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043805854, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043805854}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T05:08:32.946681+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4711416", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T14:58:38.614464+00:00", "modelId": "zstack/qwen3-4b-fake", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 25165824, "estimatedRequiredGiB": 0.028, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 25169485, "modelscopeLicense": null, "modelscopeParams": 25165435, "modelscopeTags": ["model_type:qwen3", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 25169485}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T06:46:06.785300+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712808", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T14:58:38.614275+00:00", "modelId": "OpenBMB/MiniCPM5-2B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5033557096, "estimatedRequiredGiB": 5.637, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5043805854, "modelscopeLicense": "apache-2.0", "modelscopeParams": 2516756480, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:minicpm", "custom_tag:minicpm5", "custom_tag:llama", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-calling", "custom_tag:on-device", "custom_tag:edge-ai"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5043805854}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T06:46:06.799047+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712803", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T14:58:38.614439+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a8", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14293335560, "estimatedRequiredGiB": 15.977, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 14295942065, "modelscopeLicense": "mit", "modelscopeParams": 13960238080, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 14295942065}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T06:46:06.810137+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712809", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T14:58:38.614292+00:00", "modelId": "inferencerlabs/DeepSeek-V4-Flash-MTP-DSpark-MLX", "modelProfile": {"architectures": [], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12994906286, "estimatedRequiredGiB": 14.53, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "deepseek_v4_dspark", "modelscopeFileSize": 13001289422, "modelscopeLicense": null, "modelscopeParams": 4276397927, "modelscopeTags": ["model_type:deepseek_v4_dspark", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13001289422}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T06:46:11.590454+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712813", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T14:58:38.614454+00:00", "modelId": "nanbeige/Nanbeige4.2-3B-GPTQ-Int8", "modelProfile": {"architectures": ["NanbeigeForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5192364968, "estimatedRequiredGiB": 5.827, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "nanbeige", "modelscopeFileSize": 5214325469, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4169800704, "modelscopeTags": ["license:apache-2.0", "model_type:nanbeige", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:llm", "custom_tag:nanbeige"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 5214325469}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T06:46:11.606547+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712821", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-08T14:58:38.614444+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4450260320, "estimatedRequiredGiB": 4.991, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4465562659, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "fp8", "repositoryOnDiskBytes": 4465562659}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-08T06:46:11.792714+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712824", "taskType": "text-generation", "verifyResult": null}
@@ -570,7 +568,7 @@
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-10T15:19:40.009053+00:00", "modelId": "primitive-ai/Nex-N2.5-mini-FP8", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 38136526544, "estimatedRequiredGiB": 42.65, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 38162746086, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35107181936, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:vllm", "custom_tag:compressed-tensors", "custom_tag:fp8", "custom_tag:quantized", "custom_tag:moe", "custom_tag:nex-n2.5", "custom_tag:agentic", "custom_tag:tool-calling", "custom_tag:single-gpu", "custom_tag:blackwell"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 38162746086}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:12:26.204766+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4755596", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-10T15:19:40.009045+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:12:26.209710+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4755597", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-10T15:19:40.009069+00:00", "modelId": "BAAI/AREX-Turbo", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9078620368, "estimatedRequiredGiB": 10.18, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9109259594, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4539265536, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:deep-research", "custom_tag:reasoning", "custom_tag:tool-use", "custom_tag:long-context", "custom_tag:qwen3.5", "custom_tag:dense"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "recentProfileFeedback": {"attributableFailureCount": 0, "consecutiveFailures": 0, "decisionFailureRate": 0.0, "decisionSuccessRate": 0.0, "decisionTotal": 0, "failureBreakdown": {"ambiguous_runtime": 1}, "failureCount": 1, "failureRate": 1.0, "framework": "vllm_fix_tokenizer", "lastTerminalAt": "2026-09-09T14:21:21.949441+00:00", "modelType": "qwen3_5", "pendingCount": 0, "pendingRate": 0.0, "platformFailureCount": 0, "quantizationMethod": "none", "successCount": 0, "successRate": 0.0, "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation", "total": 1, "unresolvedFailureCount": 1}, "repositoryOnDiskBytes": 9109259594}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:12:26.214438+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4755598", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T07:23:10.180003+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4755805", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-10T15:27:09.820539+00:00", "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "1566bbd6c69427ad32f7f86aeba4e39850a1a7af5c6f4ef3c9ca82f0e882f3c5", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-10T07:23:10.180003+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4755805", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T07:33:41.787375+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4756078", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T07:50:25.385741+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4756398", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-en-thinker", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 6730844445, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:reasoning", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730844445}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-10T09:52:16.471862+00:00", "targetGpu": "MetaX_c-500", "taskId": "4758203", "taskType": "text-generation", "verifyResult": null}