state: generation 18465 (cycle)
This commit is contained in:
@@ -3280,7 +3280,7 @@
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-29T17:04:24.668908+00:00",
|
||||
"generatedAt": "2026-09-29T17:16:53.495400+00:00",
|
||||
"summary": {
|
||||
"activeBlockCount": 165,
|
||||
"byGpuFramework": {
|
||||
|
||||
@@ -434,7 +434,7 @@
|
||||
}
|
||||
},
|
||||
"frameworkUpdatedAt": null,
|
||||
"generatedAt": "2026-09-29T17:13:23.070648+00:00",
|
||||
"generatedAt": "2026-09-29T17:18:26.655588+00:00",
|
||||
"gpuStats": {
|
||||
"Ascend_910-b3": {
|
||||
"available": true,
|
||||
|
||||
@@ -1,5 +1,5 @@
|
||||
{
|
||||
"catalogUpdatedAt": "2026-09-29T17:13:23.070648+00:00",
|
||||
"catalogUpdatedAt": "2026-09-29T17:18:26.655588+00:00",
|
||||
"configuredTaskTypes": [
|
||||
"text-generation"
|
||||
],
|
||||
@@ -56,7 +56,7 @@
|
||||
"time-series-forecasting"
|
||||
],
|
||||
"errors": [],
|
||||
"generatedAt": "2026-09-29T17:13:23.070648+00:00",
|
||||
"generatedAt": "2026-09-29T17:18:26.655588+00:00",
|
||||
"gpuCatalog": {
|
||||
"Ascend_910-b3": {
|
||||
"canVerify": true,
|
||||
@@ -6316,6 +6316,6 @@
|
||||
"updateTime": "2025-12-22 08:59:53"
|
||||
}
|
||||
],
|
||||
"taskTreeUpdatedAt": "2026-09-29T17:13:23.070648+00:00",
|
||||
"taskTreeUpdatedAt": "2026-09-29T17:18:26.655588+00:00",
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -1,6 +1,6 @@
|
||||
{
|
||||
"generatedAt": "2026-09-29T16:52:38.315112+00:00",
|
||||
"lastSyncTime": "2026-09-29T16:52:38.054155+00:00",
|
||||
"generatedAt": "2026-09-29T17:16:53.405682+00:00",
|
||||
"lastSyncTime": "2026-09-29T17:16:53.043636+00:00",
|
||||
"recentLimit": 300,
|
||||
"report": {
|
||||
"architectureCompatibilityBlocks": {
|
||||
@@ -3407,24 +3407,24 @@
|
||||
"decisionSuccessRate": 0.1053,
|
||||
"decisionTotal": 38,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 62,
|
||||
"ambiguous_runtime": 64,
|
||||
"context_length": 5,
|
||||
"framework_architecture_unsupported": 28,
|
||||
"memory_capacity": 1,
|
||||
"参数/模板问题": 8
|
||||
},
|
||||
"failureCount": 104,
|
||||
"failureRate": 0.963,
|
||||
"failureCount": 106,
|
||||
"failureRate": 0.9636,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 4,
|
||||
"successRate": 0.037,
|
||||
"successRate": 0.0364,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 108,
|
||||
"unresolvedFailureCount": 70
|
||||
"total": 110,
|
||||
"unresolvedFailureCount": 72
|
||||
},
|
||||
"Ascend_910-b3|vllm|text-generation": {
|
||||
"attributableFailureCount": 69,
|
||||
@@ -6130,7 +6130,7 @@
|
||||
"decisionSuccessRate": 0.0709,
|
||||
"decisionTotal": 141,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 174,
|
||||
"ambiguous_runtime": 176,
|
||||
"backend_operator": 15,
|
||||
"context_length": 5,
|
||||
"framework_architecture_unsupported": 78,
|
||||
@@ -6142,18 +6142,18 @@
|
||||
"tokenizer_compatibility": 2,
|
||||
"参数/模板问题": 29
|
||||
},
|
||||
"failureCount": 336,
|
||||
"failureRate": 0.9711,
|
||||
"failureCount": 338,
|
||||
"failureRate": 0.9713,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 2,
|
||||
"successCount": 10,
|
||||
"successRate": 0.0289,
|
||||
"total": 346,
|
||||
"unresolvedFailureCount": 203
|
||||
"successRate": 0.0287,
|
||||
"total": 348,
|
||||
"unresolvedFailureCount": 205
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-29T16:52:38.302643+00:00",
|
||||
"generatedAt": "2026-09-29T17:16:53.391898+00:00",
|
||||
"gpuSummaries": {
|
||||
"Ascend_910-b3": {
|
||||
"attributableFailureCount": 104,
|
||||
@@ -6161,7 +6161,7 @@
|
||||
"decisionSuccessRate": 0.1746,
|
||||
"decisionTotal": 126,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 128,
|
||||
"ambiguous_runtime": 130,
|
||||
"context_length": 5,
|
||||
"framework_architecture_unsupported": 93,
|
||||
"memory_capacity": 2,
|
||||
@@ -6171,15 +6171,15 @@
|
||||
"日志缺失": 3,
|
||||
"验证失败": 27
|
||||
},
|
||||
"failureCount": 301,
|
||||
"failureRate": 0.9319,
|
||||
"failureCount": 303,
|
||||
"failureRate": 0.9323,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 22,
|
||||
"successRate": 0.0681,
|
||||
"total": 323,
|
||||
"unresolvedFailureCount": 197
|
||||
"successRate": 0.0677,
|
||||
"total": 325,
|
||||
"unresolvedFailureCount": 199
|
||||
},
|
||||
"Ascend_910-b4": {
|
||||
"attributableFailureCount": 324,
|
||||
@@ -7036,9 +7036,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 4
|
||||
"ambiguous_runtime": 5
|
||||
},
|
||||
"failureCount": 4,
|
||||
"failureCount": 5,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"modelType": "phi3",
|
||||
@@ -7050,8 +7050,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 4,
|
||||
"unresolvedFailureCount": 4
|
||||
"total": 5,
|
||||
"unresolvedFailureCount": 5
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|plamo3|none": {
|
||||
"attributableFailureCount": 1,
|
||||
@@ -7471,9 +7471,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 11
|
||||
"ambiguous_runtime": 12
|
||||
},
|
||||
"failureCount": 11,
|
||||
"failureCount": 12,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"modelType": "starcoder2",
|
||||
@@ -7485,8 +7485,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 11,
|
||||
"unresolvedFailureCount": 11
|
||||
"total": 12,
|
||||
"unresolvedFailureCount": 12
|
||||
},
|
||||
"Ascend_910-b3|vllm|text-generation|bailing_hybrid|none": {
|
||||
"attributableFailureCount": 1,
|
||||
@@ -27997,9 +27997,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 2
|
||||
"ambiguous_runtime": 3
|
||||
},
|
||||
"failureCount": 2,
|
||||
"failureCount": 3,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"loadSizeLog2Bucket": 33,
|
||||
@@ -28012,8 +28012,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 2
|
||||
"total": 3,
|
||||
"unresolvedFailureCount": 3
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|plamo3|none|33": {
|
||||
"attributableFailureCount": 1,
|
||||
@@ -28595,9 +28595,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1
|
||||
"ambiguous_runtime": 2
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureCount": 2,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"loadSizeLog2Bucket": 32,
|
||||
@@ -28610,8 +28610,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b3",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 1
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 2
|
||||
},
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|starcoder2|compressed-tensors|33": {
|
||||
"attributableFailureCount": 0,
|
||||
@@ -52868,15 +52868,15 @@
|
||||
"unresolvedFailureCount": 0
|
||||
}
|
||||
},
|
||||
"terminalRecords": 17122,
|
||||
"totalRecords": 17343,
|
||||
"terminalRecords": 17124,
|
||||
"totalRecords": 17345,
|
||||
"totals": {
|
||||
"attributableFailureCount": 6136,
|
||||
"decisionFailureRate": 0.8633,
|
||||
"decisionSuccessRate": 0.1367,
|
||||
"decisionTotal": 7108,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 4449,
|
||||
"ambiguous_runtime": 4451,
|
||||
"architecture_compatibility": 212,
|
||||
"attention_backend": 3,
|
||||
"backend_operator": 128,
|
||||
@@ -52892,15 +52892,15 @@
|
||||
"日志缺失": 719,
|
||||
"验证失败": 676
|
||||
},
|
||||
"failureCount": 16150,
|
||||
"failureCount": 16152,
|
||||
"failureRate": 0.9432,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 942,
|
||||
"successCount": 972,
|
||||
"successRate": 0.0568,
|
||||
"total": 17122,
|
||||
"unresolvedFailureCount": 9072
|
||||
"total": 17124,
|
||||
"unresolvedFailureCount": 9074
|
||||
},
|
||||
"warnings": [
|
||||
"GPU Ascend_910-b3 本地统计失败率偏高(≥50%),建议重点关注。",
|
||||
@@ -52929,7 +52929,6 @@
|
||||
"组合 Kunlunxin_p-800|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b4|llamacpp|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_mrv-100|unknown|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Kunlunxin_p-800|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 MetaX_c-500|vllm|visual-multi-modal 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b4|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
@@ -52954,10 +52953,11 @@
|
||||
"组合 hygon_k100-ai|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Cambricon_mlu-370-x4|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Ascend_910-b3|vllm_tokenizer_patch|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Vastai_va16|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
"组合 Vastai_va16|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
|
||||
"组合 Kunlunxin_p-800|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。"
|
||||
]
|
||||
},
|
||||
"storageMode": "decision_state_only",
|
||||
"summarizedRecords": 17343,
|
||||
"summarizedRecords": 17345,
|
||||
"version": 1
|
||||
}
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
File diff suppressed because it is too large
Load Diff
@@ -315,7 +315,6 @@
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "submitTime": "2026-09-20T09:25:37.090994+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976380", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/bharatgenai/Param2-17B-A2.4B-Thinking", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "submitTime": "2026-09-20T09:41:34.108703+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976623", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "modelId": "neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "submitTime": "2026-09-20T11:52:06.534997+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4978249", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-2b-it-quantized.w8a16", "modelId": "RedHatAI/gemma-2-2b-it-quantized.w8a16", "submitTime": "2026-09-20T12:13:32.801809+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4978494", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-quantized.w8a8", "modelId": "RedHatAI/starcoder2-7b-quantized.w8a8", "submitTime": "2026-09-20T12:37:15.056589+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4979021", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "submitTime": "2026-09-20T13:01:42.205703+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4979280", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-iluvatar-bi-150"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-3b-quantized.w8a8", "modelId": "neuralmagic/starcoder2-3b-quantized.w8a8", "submitTime": "2026-09-20T14:38:07.337494+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4980575", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
@@ -395,7 +394,6 @@
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W8A8-Dyn-Per-Token", "modelId": "nm-testing/Meta-Llama-3-8B-Instruct-W8A8-Dyn-Per-Token", "submitTime": "2026-09-21T09:40:34.490655+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4998300", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-2b-quantized.w8a16", "modelId": "RedHatAI/gemma-2-2b-quantized.w8a16", "submitTime": "2026-09-21T09:41:35.705044+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4998344", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/blue2star/Qwen-Image-2.1-PE-I2I-ComfyUI", "modelId": "blue2star/Qwen-Image-2.1-PE-I2I-ComfyUI", "submitTime": "2026-09-21T10:34:34.644445+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4999238", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-sunrise-pt-200-x1"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-FP8", "modelId": "RedHatAI/starcoder2-7b-FP8", "submitTime": "2026-09-21T10:38:42.835320+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4999278", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "submitTime": "2026-09-21T10:38:42.837708+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4999277", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-1.2B-Thinking-MLX-bf16", "modelId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-bf16", "submitTime": "2026-09-21T10:52:01.384099+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4999444", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"}
|
||||
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-27b-it-quantized.w8a16", "modelId": "RedHatAI/gemma-2-27b-it-quantized.w8a16", "submitTime": "2026-09-21T10:52:10.635641+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4999447", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-sunrise-pt-200-x1"}
|
||||
@@ -409,7 +407,6 @@
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "modelId": "RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "submitTime": "2026-09-21T11:18:21.960935+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4999762", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-FP8", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-FP8", "submitTime": "2026-09-21T11:57:55.153765+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5000419", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w4a16", "modelId": "RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w4a16", "submitTime": "2026-09-21T12:49:01.251589+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5001057", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm", "modelAddress": "https://modelscope.cn/models/mlx-community/LFM2.5-1.2B-Thinking-6bit", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-6bit", "submitTime": "2026-09-21T13:01:28.899242+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5001186", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-ascend-910-b3"}
|
||||
{"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/inceptionai/Jais-2-8B-Chat", "modelId": "inceptionai/Jais-2-8B-Chat", "submitTime": "2026-09-21T14:18:50.561697+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "5002167", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-iluvatar-bi-150"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-2b-it-quantized.w8a8", "modelId": "RedHatAI/gemma-2-2b-it-quantized.w8a8", "submitTime": "2026-09-21T14:48:11.802242+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5002502", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/amd/Qwen3.6-35B-A3B-w4a16-llmcompressor", "modelId": "amd/Qwen3.6-35B-A3B-w4a16-llmcompressor", "submitTime": "2026-09-21T15:54:32.713857+00:00", "targetGpu": "Ascend_910-b3", "taskId": "5003279", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b3"}
|
||||
|
||||
@@ -2,24 +2,24 @@
|
||||
"agentVersion": "2026.09.22.1",
|
||||
"checksums": {
|
||||
".modelhub_state/account_capacity.json": "930f9869793c3d1aa071c53f05b7cf82a4e372f002be88361d6d6189e27c12a1",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "200939c3da6bc7fa455a1de591ece37e704c26b7a8dd9bc2c9b415357a02d00a",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "2081bcc55616d8e392d416014913fe75c92a4721e3759d9ca5f6fa4352cc122c",
|
||||
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
||||
".modelhub_state/market_intelligence.json": "d1165e81832270f99c7483e2f4131c3765de9a5457acce434b7eeb30e4758bfb",
|
||||
".modelhub_state/official_capabilities.json": "f18f2b18b841770b91acf45dab130936c70aae3893dbafc8c98cd64664aa3274",
|
||||
".modelhub_state/outcome_checkpoint.json": "f27fb0641b5e553f8e1f4e37431d632031ed403e0a27cd56018770facdf866d1",
|
||||
".modelhub_state/market_intelligence.json": "3a41e0f17e82487b3a73971771fec50a9b4ae9b2bf80148ee17b69b4e4a00f1b",
|
||||
".modelhub_state/official_capabilities.json": "a765c249c7feab1d95a45c6da8ba299126248b478762126bf37c0ac13bf51622",
|
||||
".modelhub_state/outcome_checkpoint.json": "e9ea29ae6c08191e17ef80275e11113cd14342591f3972676f713ec472ad43a9",
|
||||
".modelhub_state/queue_cleanup_latest.json": "49a85c11309f7edea696715997c353c42ca533519a99ab3a549526cccde3cad7",
|
||||
".modelhub_state/recent_outcomes.jsonl": "81e91996ffb3668bc470e02ef1be371acb37d0b6aa5850bd175907421dd7d460",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "1aa1cf9388ec44f18e81b2f70487c6e0365bbb2502fca458a6804663878aed5e",
|
||||
".modelhub_state/recovery_intents.jsonl": "08b07ba4f9795a2e54342f78b15249e38650304fb131f2392cc368bc1d0c05b6",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "34e6070c0f9ea5951606f9b1696d7aef8304876932138ddfbdc3fdc312ee5c78",
|
||||
".modelhub_state/recovery_intents.jsonl": "9bc542594a7976f342c33ed1594c1aff3be13f3ec399b5d5376815228d60b953",
|
||||
".modelhub_state/routing_intelligence.json": "895431dc6ea0fddce4c8a244da61cbaa7bdfb36cfdbd17e79f8e50e130dca87a",
|
||||
".modelhub_state/submission_exclusions.jsonl": "1e62f6ba5bd2ba5f2f360bc134b44f7b69190230dd0e274919a87a73ea66761c",
|
||||
".modelhub_state/worker_crashes.jsonl": "05bcde79076023fd20785b8270312f776430b95a2d4322ee53ef24e354c39f06",
|
||||
"ledger/submissions.jsonl": "5ebc62dc59588526d925e4bdc7fb66863ed801aa041e0f05c390e0f2ce80ba2f",
|
||||
"outcomes/submissions.jsonl": "d2df412060df36da65d698d3e8d489bc259ab438054fe9e9997bdf736b836c60"
|
||||
"ledger/submissions.jsonl": "4e38a4e04b45021e3a767a110b76ae47d9635675eea4622c2e146306e2feb9a1",
|
||||
"outcomes/submissions.jsonl": "1d385e9d0e0688d75933fd8c1df41a6174887206c5fcc25feb9642473b2f36e3"
|
||||
},
|
||||
"generation": 18464,
|
||||
"generation": 18465,
|
||||
"phase": "cycle",
|
||||
"schemaVersion": 1,
|
||||
"updatedAt": "2026-09-29T17:15:50.143089+00:00",
|
||||
"updatedAt": "2026-09-29T17:18:34.595561+00:00",
|
||||
"writerId": "9d958b95ce4146c9939da0f14b9201ff"
|
||||
}
|
||||
|
||||
@@ -312,7 +312,6 @@
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:21:49.901858+00:00", "modelId": "inceptionai/Jais-2-8B-Chat", "modelProfile": {"architectures": ["Jais2ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16180861064, "estimatedRequiredGiB": 18.097, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "jais2", "modelscopeFileSize": 16193187697, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8090401280, "modelscopeTags": ["license:apache-2.0", "model_type:jais2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16193187697}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:10:03.058495+00:00", "targetGpu": "MetaX_c-500", "taskId": "4970512", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:21:49.901772+00:00", "modelId": "mlx-community/LFM2.5-1.2B-Thinking-8bit", "modelProfile": {"architectures": ["Lfm2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 1243645809, "estimatedRequiredGiB": 1.395, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "lfm2", "modelscopeFileSize": 1248409935, "modelscopeLicense": "other", "modelscopeParams": 329251584, "modelscopeTags": ["license:other", "model_type:lfm2", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:liquid", "custom_tag:lfm2.5", "custom_tag:edge", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1248409935}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:11:06.683086+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970527", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:21:49.901659+00:00", "modelId": "RedHatAI/Meta-Llama-3.1-70B-Instruct-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 72670025872, "estimatedRequiredGiB": 81.225, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 72679340864, "modelscopeLicense": "llama3.1", "modelscopeParams": 70553706496, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int88", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 72679340864}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:14:34.788367+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4970568", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:21:49.901883+00:00", "modelId": "neuralmagic/starcoder2-7b-quantized.w8a16", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7857298896, "estimatedRequiredGiB": 8.785, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7860673858, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 7400416256, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7860673858}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:14:34.922099+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970574", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T12:21:49.902090+00:00", "modelId": "RedHatAI/Qwen2-1.5B-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2245331488, "estimatedRequiredGiB": 2.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 2256941144, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1777088000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2256941144}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:14:34.856488+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4970570", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T12:21:49.901918+00:00", "modelId": "RedHatAI/gemma-2-27b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 30775446656, "estimatedRequiredGiB": 34.419, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 30797371689, "modelscopeLicense": "gemma", "modelscopeParams": 28406776320, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 30797371689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:14:42.093442+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4970600", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T12:21:49.901870+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T04:15:36.946953+00:00", "targetGpu": "Vastai_va16", "taskId": "4970616", "taskType": "text-generation", "verifyResult": null}
|
||||
@@ -328,7 +327,6 @@
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T17:39:36.352398+00:00", "modelId": "RedHatAI/Llama-3.2-3B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4395007696, "estimatedRequiredGiB": 4.922, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 4404161034, "modelscopeLicense": "llama3.2", "modelscopeParams": 3606752256, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4404161034}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T09:25:37.077913+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976379", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T17:39:36.352379+00:00", "modelId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4020365960, "estimatedRequiredGiB": 4.496, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 4022810352, "modelscopeLicense": "mit", "modelscopeParams": 3821079552, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4022810352}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T09:25:37.090994+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976380", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T17:54:13.170240+00:00", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "modelProfile": {"architectures": ["Param2MoEForCausalLM"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34302758832, "estimatedRequiredGiB": 38.353, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "param2moe", "modelscopeFileSize": 34317809830, "modelscopeLicense": null, "modelscopeParams": 17151125376, "modelscopeTags": ["model_type:param2moe", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:mixture-of-experts", "custom_tag:multilingual", "custom_tag:indian-languages"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 34317809830}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T09:41:34.108703+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976623", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T19:33:34.552425+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 14293356304, "estimatedRequiredGiB": 15.977, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 14295851911, "modelscopeLicense": "mit", "modelscopeParams": 13960238080, "modelscopeTags": ["license:mit", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 14295851911}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T11:19:51.641136+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4977829", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T19:56:29.962409+00:00", "modelId": "neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7686303632, "estimatedRequiredGiB": 8.592, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 7688299919, "modelscopeLicense": "llama2", "modelscopeParams": 13960238080, "modelscopeTags": ["license:llama2", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7688299919}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T11:52:06.534997+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4978249", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:16:30.961235+00:00", "modelId": "RedHatAI/gemma-2-2b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4385544296, "estimatedRequiredGiB": 4.926, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 4407367933, "modelscopeLicense": "gemma", "modelscopeParams": 3204165888, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4407367933}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T12:13:32.801809+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4978494", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:49:19.683309+00:00", "modelId": "RedHatAI/starcoder2-7b-quantized.w8a8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7857274344, "estimatedRequiredGiB": 8.785, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7860631926, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 7400416256, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7860631926}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T12:37:15.056589+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4979021", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user