state: generation 4848 (cycle)

This commit is contained in:
2026-09-10 00:05:29 +00:00
parent 8f5ee0299a
commit 1daca9712a
10 changed files with 1971 additions and 1977 deletions

View File

@@ -1312,7 +1312,7 @@
"taskType": "text-generation"
}
},
"generatedAt": "2026-09-10T00:02:39.848956+00:00",
"generatedAt": "2026-09-10T00:05:29.186799+00:00",
"summary": {
"activeBlockCount": 65,
"byGpuFramework": {

View File

@@ -19,7 +19,7 @@
"10": {
"complete": false,
"lastError": "ModelHubAPIError: 系统错误",
"listingErrors": 344,
"listingErrors": 345,
"nextPage": 1,
"recordsScanned": 0,
"uniqueRecords": 0
@@ -102,12 +102,12 @@
"cutoffAt": "2026-09-04T03:55:51.365685+00:00",
"failureLogsInspected": 0,
"mode": "incremental_decision_only",
"nextAccountIndex": 10,
"nextAccountIndex": 11,
"recordsScanned": 0,
"seenTaskIds": [],
"startedAt": "2026-09-04T03:55:51.365685+00:00",
"terminalRecords": 0,
"uniqueRecords": 0,
"updatedAt": "2026-09-10T00:02:39.826633+00:00",
"updatedAt": "2026-09-10T00:05:29.162194+00:00",
"version": 1
}

View File

@@ -416,7 +416,7 @@
}
},
"frameworkUpdatedAt": null,
"generatedAt": "2026-09-10T00:01:17.038580+00:00",
"generatedAt": "2026-09-10T00:03:48.111564+00:00",
"gpuStats": {
"Ascend_910-b3": {
"available": true,

File diff suppressed because it is too large Load Diff

View File

@@ -1,6 +1,6 @@
{
"generatedAt": "2026-09-09T23:55:17.575624+00:00",
"lastSyncTime": "2026-09-09T23:55:17.200912+00:00",
"generatedAt": "2026-09-10T00:03:40.752831+00:00",
"lastSyncTime": "2026-09-10T00:03:40.507646+00:00",
"recentLimit": 300,
"report": {
"architectureCompatibilityBlocks": {
@@ -1906,19 +1906,19 @@
"unresolvedFailureCount": 56
},
"MetaX_c-500|vllm|text-generation": {
"attributableFailureCount": 36,
"attributableFailureCount": 37,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 36,
"decisionTotal": 37,
"failureBreakdown": {
"ambiguous_runtime": 7,
"backend_operator": 17,
"backend_operator": 18,
"framework_architecture_unsupported": 14,
"memory_capacity": 1,
"model_load": 4,
"参数/模板问题": 9
},
"failureCount": 52,
"failureCount": 53,
"failureRate": 1.0,
"framework": "vllm",
"pendingCount": 0,
@@ -1928,7 +1928,7 @@
"successRate": 0.0,
"targetGpu": "MetaX_c-500",
"taskType": "text-generation",
"total": 52,
"total": 53,
"unresolvedFailureCount": 16
},
"Mthreads_s4000|llamacpp|text-generation": {
@@ -2315,13 +2315,13 @@
"unresolvedFailureCount": 819
},
"vllm": {
"attributableFailureCount": 225,
"attributableFailureCount": 226,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 225,
"decisionTotal": 226,
"failureBreakdown": {
"ambiguous_runtime": 120,
"backend_operator": 20,
"backend_operator": 21,
"framework_architecture_unsupported": 168,
"memory_capacity": 6,
"model_load": 18,
@@ -2331,14 +2331,14 @@
"tokenizer_compatibility": 3,
"参数/模板问题": 21
},
"failureCount": 367,
"failureCount": 368,
"failureRate": 1.0,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 1,
"successCount": 0,
"successRate": 0.0,
"total": 367,
"total": 368,
"unresolvedFailureCount": 141
},
"vllm-mlu": {
@@ -2442,7 +2442,7 @@
"unresolvedFailureCount": 1
}
},
"generatedAt": "2026-09-09T23:55:17.571160+00:00",
"generatedAt": "2026-09-10T00:03:40.748614+00:00",
"gpuSummaries": {
"Ascend_910-b3": {
"attributableFailureCount": 26,
@@ -2649,27 +2649,27 @@
"unresolvedFailureCount": 72
},
"MetaX_c-500": {
"attributableFailureCount": 36,
"decisionFailureRate": 0.8,
"decisionSuccessRate": 0.2,
"decisionTotal": 45,
"attributableFailureCount": 37,
"decisionFailureRate": 0.8043,
"decisionSuccessRate": 0.1957,
"decisionTotal": 46,
"failureBreakdown": {
"ambiguous_runtime": 7,
"backend_operator": 17,
"backend_operator": 18,
"framework_architecture_unsupported": 14,
"memory_capacity": 1,
"model_load": 4,
"参数/模板问题": 17,
"验证失败": 48
},
"failureCount": 108,
"failureRate": 0.9231,
"failureCount": 109,
"failureRate": 0.9237,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"successCount": 9,
"successRate": 0.0769,
"total": 117,
"successRate": 0.0763,
"total": 118,
"unresolvedFailureCount": 72
},
"Mthreads_s4000": {
@@ -3614,14 +3614,14 @@
"unresolvedFailureCount": 0
},
"MetaX_c-500|vllm|text-generation|starcoder2|compressed-tensors": {
"attributableFailureCount": 5,
"attributableFailureCount": 6,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 5,
"decisionTotal": 6,
"failureBreakdown": {
"backend_operator": 5
"backend_operator": 6
},
"failureCount": 5,
"failureCount": 6,
"failureRate": 1.0,
"framework": "vllm",
"modelType": "starcoder2",
@@ -3633,7 +3633,7 @@
"successRate": 0.0,
"targetGpu": "MetaX_c-500",
"taskType": "text-generation",
"total": 5,
"total": 6,
"unresolvedFailureCount": 0
},
"Mthreads_s4000|vllm|text-generation|glm_ocr|fp8": {
@@ -6366,14 +6366,14 @@
"unresolvedFailureCount": 0
},
"MetaX_c-500|vllm|text-generation|starcoder2|compressed-tensors|34": {
"attributableFailureCount": 1,
"attributableFailureCount": 2,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"decisionTotal": 2,
"failureBreakdown": {
"backend_operator": 1
"backend_operator": 2
},
"failureCount": 1,
"failureCount": 2,
"failureRate": 1.0,
"framework": "vllm",
"loadSizeLog2Bucket": 34,
@@ -6386,7 +6386,7 @@
"successRate": 0.0,
"targetGpu": "MetaX_c-500",
"taskType": "text-generation",
"total": 1,
"total": 2,
"unresolvedFailureCount": 0
},
"Mthreads_s4000|vllm|text-generation|glm_ocr|fp8|30": {
@@ -6678,16 +6678,16 @@
"unresolvedFailureCount": 0
}
},
"terminalRecords": 1459,
"totalRecords": 1543,
"terminalRecords": 1460,
"totalRecords": 1544,
"totals": {
"attributableFailureCount": 388,
"decisionFailureRate": 0.8899,
"decisionSuccessRate": 0.1101,
"decisionTotal": 436,
"attributableFailureCount": 389,
"decisionFailureRate": 0.8902,
"decisionSuccessRate": 0.1098,
"decisionTotal": 437,
"failureBreakdown": {
"ambiguous_runtime": 254,
"backend_operator": 23,
"backend_operator": 24,
"framework_architecture_unsupported": 266,
"memory_capacity": 11,
"model_load": 45,
@@ -6698,21 +6698,21 @@
"参数/模板问题": 94,
"验证失败": 673
},
"failureCount": 1411,
"failureCount": 1412,
"failureRate": 0.9671,
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 2,
"successCount": 48,
"successRate": 0.0329,
"total": 1459,
"total": 1460,
"unresolvedFailureCount": 1021
},
"warnings": [
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_mrv-100 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Sunrise_pt-200-x1 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
@@ -6720,7 +6720,7 @@
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Mthreads_s4000 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Sunrise_pt-200-x1 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Iluvatar_mrv-100|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Cambricon_mlu-370-x8|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -6741,6 +6741,6 @@
]
},
"storageMode": "decision_state_only",
"summarizedRecords": 1543,
"summarizedRecords": 1544,
"version": 1
}

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -121,7 +121,6 @@
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-Non-Uniform-compressed-tensors", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-Non-Uniform-compressed-tensors", "submitTime": "2026-09-05T09:38:43.506267+00:00", "targetGpu": "MetaX_c-500", "taskId": "4645748", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-metax-c-500"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-FP8-K-V", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-FP8-K-V", "submitTime": "2026-09-05T09:38:43.485599+00:00", "targetGpu": "MetaX_c-500", "taskId": "4645738", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-metax-c-500"}
{"framework": "vllm_0_17_0_corex_4_4_0", "modelAddress": "https://modelscope.cn/models/Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B", "modelId": "Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B", "submitTime": "2026-09-05T09:43:33.758884+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4645816", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-0-17-0-corex-4-4-0-iluvatar-bi-150"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W4A16-compressed-tensors-test", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-compressed-tensors-test", "submitTime": "2026-09-05T10:33:51.489696+00:00", "targetGpu": "MetaX_c-500", "taskId": "4646504", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-metax-c-500"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/zstack/qwen3-4b-fake", "modelId": "zstack/qwen3-4b-fake", "submitTime": "2026-09-05T10:33:51.484756+00:00", "targetGpu": "MetaX_c-500", "taskId": "4646505", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-metax-c-500"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "submitTime": "2026-09-05T11:12:23.284350+00:00", "targetGpu": "MetaX_c-500", "taskId": "4647279", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-metax-c-500"}
{"framework": "llamacpp", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "submitTime": "2026-09-05T12:03:37.084154+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4648146", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-llamacpp-hygon-k100-ai"}
@@ -207,10 +206,8 @@
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/siliconflow/gpt-oss-20b-FP8", "modelId": "siliconflow/gpt-oss-20b-FP8", "submitTime": "2026-09-06T14:32:32.785339+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4668538", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/RedHatAI/Meta-Llama-3.1-70B-Instruct-FP8-dynamic", "modelId": "RedHatAI/Meta-Llama-3.1-70B-Instruct-FP8-dynamic", "submitTime": "2026-09-06T14:33:50.986624+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4668642", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-3b-FP8", "modelId": "RedHatAI/starcoder2-3b-FP8", "submitTime": "2026-09-06T14:34:13.216402+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4668650", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-FP8-K-V", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-FP8-K-V", "submitTime": "2026-09-06T14:52:34.105196+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4668969", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/zstack/qwen3-4b-fake", "modelId": "zstack/qwen3-4b-fake", "submitTime": "2026-09-06T15:05:58.785747+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669358", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/cyankiwi/Ornith-1.5-9B-AWQ-FP8", "modelId": "cyankiwi/Ornith-1.5-9B-AWQ-FP8", "submitTime": "2026-09-06T15:07:54.384483+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669411", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "submitTime": "2026-09-06T15:28:09.693682+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669866", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/RWKV/RWKV7-7.2B-20260805", "modelId": "RWKV/RWKV7-7.2B-20260805", "submitTime": "2026-09-06T16:48:59.678173+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4672290", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/RWKV/RWKV7-1.5B-20260805", "modelId": "RWKV/RWKV7-1.5B-20260805", "submitTime": "2026-09-06T18:34:36.628828+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4674757", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B", "modelId": "Jackrong/DeepSeek-V4-Pro-Qwen3.5-9B", "submitTime": "2026-09-06T19:48:32.773868+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4676609", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"}
@@ -389,7 +386,6 @@
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/whq1111/M2RL-RL_IF", "modelId": "whq1111/M2RL-RL_IF", "submitTime": "2026-09-08T22:49:56.923972+00:00", "targetGpu": "Biren_166m", "taskId": "4729274", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-biren-166m"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/TokenRhythm/NeoHorse-1-4B", "modelId": "TokenRhythm/NeoHorse-1-4B", "submitTime": "2026-09-08T22:49:56.925801+00:00", "targetGpu": "Biren_166m", "taskId": "4729275", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-biren-166m"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/OpenBMB/MiniCPM5-2B", "modelId": "OpenBMB/MiniCPM5-2B", "submitTime": "2026-09-08T23:02:45.687861+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4729507", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W4A16-compressed-tensors-test", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-compressed-tensors-test", "submitTime": "2026-09-08T23:02:45.685300+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4729504", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/whq1111/M2RL-SFT", "modelId": "whq1111/M2RL-SFT", "submitTime": "2026-09-08T23:02:45.664222+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4729503", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/whq1111/M2RL-RL_Math", "modelId": "whq1111/M2RL-RL_Math", "submitTime": "2026-09-08T23:02:45.690523+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4729502", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/whq1111/M2RL-RL_Science", "modelId": "whq1111/M2RL-RL_Science", "submitTime": "2026-09-08T23:02:45.686402+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4729506", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
@@ -523,12 +519,15 @@
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/FlagRelease/DeepSeek-R1-Distill-Qwen-1.5B-mthreads-FlagOS", "modelId": "FlagRelease/DeepSeek-R1-Distill-Qwen-1.5B-mthreads-FlagOS", "submitTime": "2026-09-09T21:08:32.656390+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746454", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/nm-testing/llama7b-one-shot-2_4-w4a16-marlin24-t-alt", "modelId": "nm-testing/llama7b-one-shot-2_4-w4a16-marlin24-t-alt", "submitTime": "2026-09-06T03:50:43.687587+00:00", "targetGpu": "MetaX_c-500", "taskId": "4660391", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-metax-c-500"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W8A8-FP8-Channelwise-compressed-tensors", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W8A8-FP8-Channelwise-compressed-tensors", "submitTime": "2026-09-06T14:32:40.056825+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4668539", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-FP8-K-V", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-FP8-K-V", "submitTime": "2026-09-06T14:52:34.105196+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4668969", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/neuralmagic/SmolLM-1.7B-Instruct-quantized.w8a16", "modelId": "neuralmagic/SmolLM-1.7B-Instruct-quantized.w8a16", "submitTime": "2026-09-06T15:28:09.689378+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669860", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-FP8", "modelId": "neuralmagic/starcoder2-15b-FP8", "submitTime": "2026-09-06T15:28:09.687642+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669862", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-9b-it-quantized.w8a16", "modelId": "neuralmagic/gemma-2-9b-it-quantized.w8a16", "submitTime": "2026-09-06T15:28:09.690715+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669864", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/neuralmagic/SmolLM-360M-Instruct-quantized.w8a8", "modelId": "neuralmagic/SmolLM-360M-Instruct-quantized.w8a8", "submitTime": "2026-09-06T15:28:09.692486+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669865", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-ACTORDER-compressed-tensors-test", "submitTime": "2026-09-06T15:28:09.693682+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669866", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm-mlu", "modelAddress": "https://modelscope.cn/models/nm-testing/nonuniform", "modelId": "nm-testing/nonuniform", "submitTime": "2026-09-08T06:46:06.780642+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712802", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mlu-cambricon-mlu-370-x8"}
{"framework": "vllm-mlu", "modelAddress": "https://modelscope.cn/models/RWKV/RWKV7-7.2B-20260805", "modelId": "RWKV/RWKV7-7.2B-20260805", "submitTime": "2026-09-08T06:46:11.589203+00:00", "targetGpu": "Cambricon_mlu-370-x8", "taskId": "4712819", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-mlu-cambricon-mlu-370-x8"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W4A16-compressed-tensors-test", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W4A16-compressed-tensors-test", "submitTime": "2026-09-08T23:02:45.685300+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4729504", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "submitTime": "2026-09-08T23:02:51.847566+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4729520", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}
{"framework": "vllm_0_17_0_corex_4_4_0", "modelAddress": "https://modelscope.cn/models/RWKV/RWKV7-1.5B-20260805", "modelId": "RWKV/RWKV7-1.5B-20260805", "submitTime": "2026-09-05T13:47:55.466466+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4649912", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-0-17-0-corex-4-4-0-iluvatar-bi-150"}
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-quantized.w8a16", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a16", "submitTime": "2026-09-06T15:28:09.684321+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4669863", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"}

View File

@@ -1,24 +1,24 @@
{
"agentVersion": "2026.09.04.4",
"checksums": {
".modelhub_state/architecture_compatibility_blacklist.json": "ba9eb22b50a3f376bb6e1ac430a284b8c1531720d1395cf4a225e6b44dbfdc50",
".modelhub_state/architecture_history_backfill.json": "0671220584ed8ce0d3542e53a4a0e4e11b42cb8a4c0cc8ce2f51e8ce5d82844b",
".modelhub_state/market_intelligence.json": "cc3b6bf9c122911ecbf675ebecbaba68fe8ebae2eddfd3084b6d55e497eeea9d",
".modelhub_state/official_capabilities.json": "5bc2bd4fd9908b58a638f8a22c13941225b54701288fb4f19b5830cee8199a39",
".modelhub_state/outcome_checkpoint.json": "c050cbc76893bef8035b9111cf5788dd6690c6d4a4d5b87ec4a2d354b15c82ce",
".modelhub_state/architecture_compatibility_blacklist.json": "db03cdbce47595f53164e57b5f7133af78d21a843342f595b9c8d444c0e271b3",
".modelhub_state/architecture_history_backfill.json": "da150227294f04e71420f5f921e2330aefe1a479e8bb87ae691c7af5daa46b05",
".modelhub_state/market_intelligence.json": "d632abc1868ecb0d91771874c5ec4c88682dadd522e5e8a504ad96862491cae7",
".modelhub_state/official_capabilities.json": "da26d22899eff719f9cf5129023526e07e1db7481fd35c93aa9a0a9b33816ed7",
".modelhub_state/outcome_checkpoint.json": "4d8117315ea31b55885f7951199b5437c9a6c21d9862af233ceaba86acf6bb40",
".modelhub_state/queue_cleanup_latest.json": "9a11c448074d1c105fe90e48fb81cfc27bce5e7e08d76397232f05c4c9432a32",
".modelhub_state/recent_outcomes.jsonl": "4df4d35d68f5d08217f54c911ccd6a963ed8fbe9fc0bfbb30db65737b0260302",
".modelhub_state/recovery_active_tasks.jsonl": "c025690630ea1ef28d9a1e811df70f511ee6af9727171ecbc82101c88e035827",
".modelhub_state/recovery_intents.jsonl": "69e38cb80ce8f5700fb0d88fea6379cf4eff20862389ee310a2d1005a51e5152",
".modelhub_state/recovery_active_tasks.jsonl": "4a8fcc6367c2aad229909ce521556aa764921b3e29926f3319234e33a83f57ae",
".modelhub_state/recovery_intents.jsonl": "86508796e4769b8f0947988437ed4cf188931bd81cac8b19c46f1e763438a3d9",
".modelhub_state/routing_intelligence.json": "4effdc96741bb83e9319ea4665226a7e2442dc48c76af38dae47ff1343260b42",
".modelhub_state/submission_exclusions.jsonl": "363b21d9f5526f9b072e4f5ff55cc1ff126e844da64fb08a7b6e4d09b2559cdb",
".modelhub_state/worker_crashes.jsonl": "d1c3f447ce0377ed5e820c2b5ff730d29f84a96af4fd14332b138f78c0ddde76",
"ledger/submissions.jsonl": "264925874253098e21a982f3db1058bfa23af60b2724a211f3e3b6e07b2f1427",
"outcomes/submissions.jsonl": "69297aa6d5f250bb1b98daab36f5a10014d9aab753dcde0cd907400fd1894396"
"ledger/submissions.jsonl": "024b5fef68d5c10478f073674c5f7caac3d29e260c1a4743a8a41fb38a5f0de0",
"outcomes/submissions.jsonl": "8763b29aa449714ee872a2ed23fbd9845d878204e54ebba29ea2801e2cb45acc"
},
"generation": 4847,
"generation": 4848,
"phase": "cycle",
"schemaVersion": 1,
"updatedAt": "2026-09-10T00:02:39.915947+00:00",
"updatedAt": "2026-09-10T00:05:29.712210+00:00",
"writerId": "58ff5dc2a2d64b119ad1e4560700f977"
}

View File

@@ -47,7 +47,6 @@
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-04T16:50:26.817445+00:00", "modelId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15338408461, "estimatedRequiredGiB": 17.168, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 15362064483, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7669052656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:exl3", "custom_tag:exllamav3", "custom_tag:quantization", "custom_tag:long-context", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "exl3", "repositoryOnDiskBytes": 15362064483}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T08:44:58.161709+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4623226", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-04T19:52:10.096489+00:00", "modelId": "neuralmagic/SmolLM-360M-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "5dc29cb4673042e118df9721ee1fc20ce84e880992fba859071f1332bb52342f", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 504052632, "estimatedRequiredGiB": 0.567, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:29"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 507443106, "modelscopeLicense": "apache-2.0", "modelscopeParams": 409007040, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 507443106}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T11:51:30.613200+00:00", "targetGpu": "MetaX_c-500", "taskId": "4625822", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-04T20:39:27.718139+00:00", "modelId": "ornith-ai/Ornith-1.5-9B-NVFP4", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8818103892, "estimatedRequiredGiB": 9.883, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:8"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 8842796317, "modelscopeLicense": "mit", "modelscopeParams": 6728625904, "modelscopeTags": ["license:mit", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "modelopt", "repositoryOnDiskBytes": 8842796317}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T12:34:49.403585+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4626360", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-04T20:54:37.415689+00:00", "modelId": "neuralmagic/starcoder2-15b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17785269848, "estimatedRequiredGiB": 19.88, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:29"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 17788663867, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 15957889024, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 17788663867}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T12:51:02.061439+00:00", "targetGpu": "MetaX_c-500", "taskId": "4626559", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-04T20:54:37.415713+00:00", "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4395007696, "estimatedRequiredGiB": 4.922, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:29"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 4404161088, "modelscopeLicense": "llama3.2", "modelscopeParams": 3606752256, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4404161088}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T12:51:49.077105+00:00", "targetGpu": "MetaX_c-500", "taskId": "4626568", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-04T21:36:16.816860+00:00", "modelId": "siliconflow/gpt-oss-20b-FP8", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 22109595640, "estimatedRequiredGiB": 24.741, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:29"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 22137550529, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 20921584848, "modelscopeTags": ["license:Apache License 2.0", "model_type:gpt_oss", "library:safetensors", "library:other", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "mxfp4", "repositoryOnDiskBytes": 22137550529}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T13:34:49.450482+00:00", "targetGpu": "MetaX_c-500", "taskId": "4627250", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-04T21:46:36.195225+00:00", "modelId": "sapientinc/HRM-Text-1B", "modelProfile": {"architectures": ["HrmTextForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2365606568, "estimatedRequiredGiB": 2.651, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "hrm_text", "modelscopeFileSize": 2371660433, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1182795264, "modelscopeTags": ["license:apache-2.0", "model_type:hrm_text", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:hrm", "custom_tag:hierarchical-reasoning", "custom_tag:prefix-lm", "custom_tag:pre-alignment", "custom_tag:non-chat", "custom_tag:non-instruction-tuned"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 2371660433}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T13:43:30.537609+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4627500", "taskType": "text-generation", "verifyResult": null}