state: generation 11335 (intent)
This commit is contained in:
@@ -1720,6 +1720,27 @@
|
||||
"targetGpu": "Iluvatar_bi-150",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"kunlunxin_p-800|vllm_tokenizer_patch|text-generation|architectures:spark2_5forcausallm": {
|
||||
"architectureSignature": "architectures:spark2_5forcausallm",
|
||||
"architectures": [
|
||||
"spark2_5forcausallm"
|
||||
],
|
||||
"evidenceCount": 1,
|
||||
"expiresAt": "2026-10-20T03:25:31.991647+00:00",
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"latestFailureAt": "2026-09-20T03:25:31.991647+00:00",
|
||||
"latestSuccessfulAt": null,
|
||||
"matchType": "architectures",
|
||||
"modelType": null,
|
||||
"sourceModelIds": [
|
||||
"XHToken/Spark-X2.5-1.7B"
|
||||
],
|
||||
"sourceTaskIds": [
|
||||
"4969636"
|
||||
],
|
||||
"targetGpu": "Kunlunxin_p-800",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"kunlunxin_p-800|vllm_tokenizer_patch|text-generation|model_type:nemotron_h": {
|
||||
"architectureSignature": "model_type:nemotron_h",
|
||||
"architectures": [],
|
||||
@@ -2635,9 +2656,9 @@
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-21T08:22:05.179335+00:00",
|
||||
"generatedAt": "2026-09-21T08:25:33.731333+00:00",
|
||||
"summary": {
|
||||
"activeBlockCount": 132,
|
||||
"activeBlockCount": 133,
|
||||
"byGpuFramework": {
|
||||
"Ascend_910-b3|vllm": 13,
|
||||
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
||||
@@ -2652,7 +2673,7 @@
|
||||
"Iluvatar_bi-150|vllm": 7,
|
||||
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0": 3,
|
||||
"Kunlunxin_p-800|vllm": 1,
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch": 4,
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch": 5,
|
||||
"MetaX_c-500|vllm": 7,
|
||||
"Mthreads_s4000|vllm": 9,
|
||||
"Sunrise_pt-200-x1|vllm": 3,
|
||||
|
||||
@@ -425,7 +425,7 @@
|
||||
}
|
||||
},
|
||||
"frameworkUpdatedAt": null,
|
||||
"generatedAt": "2026-09-21T08:24:30.353529+00:00",
|
||||
"generatedAt": "2026-09-21T08:25:41.976766+00:00",
|
||||
"gpuStats": {
|
||||
"Ascend_910-b3": {
|
||||
"available": true,
|
||||
|
||||
File diff suppressed because it is too large
Load Diff
@@ -1,6 +1,6 @@
|
||||
{
|
||||
"generatedAt": "2026-09-21T08:22:05.105619+00:00",
|
||||
"lastSyncTime": "2026-09-21T08:22:03.151517+00:00",
|
||||
"generatedAt": "2026-09-21T08:25:41.905272+00:00",
|
||||
"lastSyncTime": "2026-09-21T08:25:41.844650+00:00",
|
||||
"recentLimit": 300,
|
||||
"report": {
|
||||
"architectureCompatibilityBlocks": {
|
||||
@@ -1724,6 +1724,27 @@
|
||||
"targetGpu": "Iluvatar_bi-150",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"kunlunxin_p-800|vllm_tokenizer_patch|text-generation|architectures:spark2_5forcausallm": {
|
||||
"architectureSignature": "architectures:spark2_5forcausallm",
|
||||
"architectures": [
|
||||
"spark2_5forcausallm"
|
||||
],
|
||||
"evidenceCount": 1,
|
||||
"expiresAt": "2026-10-20T03:25:31.991647+00:00",
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"latestFailureAt": "2026-09-20T03:25:31.991647+00:00",
|
||||
"latestSuccessfulAt": null,
|
||||
"matchType": "architectures",
|
||||
"modelType": null,
|
||||
"sourceModelIds": [
|
||||
"XHToken/Spark-X2.5-1.7B"
|
||||
],
|
||||
"sourceTaskIds": [
|
||||
"4969636"
|
||||
],
|
||||
"targetGpu": "Kunlunxin_p-800",
|
||||
"taskType": "text-generation"
|
||||
},
|
||||
"kunlunxin_p-800|vllm_tokenizer_patch|text-generation|model_type:nemotron_h": {
|
||||
"architectureSignature": "model_type:nemotron_h",
|
||||
"architectures": [],
|
||||
@@ -2640,7 +2661,7 @@
|
||||
}
|
||||
},
|
||||
"architectureCompatibilitySummary": {
|
||||
"activeBlockCount": 132,
|
||||
"activeBlockCount": 133,
|
||||
"byGpuFramework": {
|
||||
"Ascend_910-b3|vllm": 13,
|
||||
"Ascend_910-b3|vllm_tokenizer_patch": 5,
|
||||
@@ -2655,7 +2676,7 @@
|
||||
"Iluvatar_bi-150|vllm": 7,
|
||||
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0": 3,
|
||||
"Kunlunxin_p-800|vllm": 1,
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch": 4,
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch": 5,
|
||||
"MetaX_c-500|vllm": 7,
|
||||
"Mthreads_s4000|vllm": 9,
|
||||
"Sunrise_pt-200-x1|vllm": 3,
|
||||
@@ -4280,16 +4301,16 @@
|
||||
"unresolvedFailureCount": 93
|
||||
},
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation": {
|
||||
"attributableFailureCount": 10,
|
||||
"attributableFailureCount": 11,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 10,
|
||||
"decisionTotal": 11,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1,
|
||||
"framework_architecture_unsupported": 10,
|
||||
"framework_architecture_unsupported": 11,
|
||||
"参数/模板问题": 11
|
||||
},
|
||||
"failureCount": 22,
|
||||
"failureCount": 23,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"pendingCount": 0,
|
||||
@@ -4299,7 +4320,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Kunlunxin_p-800",
|
||||
"taskType": "text-generation",
|
||||
"total": 22,
|
||||
"total": 23,
|
||||
"unresolvedFailureCount": 12
|
||||
},
|
||||
"Kunlunxin_p-800|vllm|text-generation": {
|
||||
@@ -5424,32 +5445,32 @@
|
||||
"unresolvedFailureCount": 148
|
||||
},
|
||||
"vllm_tokenizer_patch": {
|
||||
"attributableFailureCount": 59,
|
||||
"decisionFailureRate": 0.9672,
|
||||
"decisionSuccessRate": 0.0328,
|
||||
"decisionTotal": 61,
|
||||
"attributableFailureCount": 60,
|
||||
"decisionFailureRate": 0.9677,
|
||||
"decisionSuccessRate": 0.0323,
|
||||
"decisionTotal": 62,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 51,
|
||||
"context_length": 1,
|
||||
"framework_architecture_unsupported": 54,
|
||||
"framework_architecture_unsupported": 55,
|
||||
"memory_capacity": 1,
|
||||
"model_load": 2,
|
||||
"platform_infrastructure": 1,
|
||||
"tokenizer_compatibility": 1,
|
||||
"参数/模板问题": 19
|
||||
},
|
||||
"failureCount": 130,
|
||||
"failureRate": 0.9848,
|
||||
"failureCount": 131,
|
||||
"failureRate": 0.985,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 1,
|
||||
"successCount": 2,
|
||||
"successRate": 0.0152,
|
||||
"total": 132,
|
||||
"successRate": 0.015,
|
||||
"total": 133,
|
||||
"unresolvedFailureCount": 70
|
||||
}
|
||||
},
|
||||
"generatedAt": "2026-09-21T08:22:05.094989+00:00",
|
||||
"generatedAt": "2026-09-21T08:25:41.893418+00:00",
|
||||
"gpuSummaries": {
|
||||
"Ascend_910-b3": {
|
||||
"attributableFailureCount": 91,
|
||||
@@ -5674,26 +5695,26 @@
|
||||
"unresolvedFailureCount": 714
|
||||
},
|
||||
"Kunlunxin_p-800": {
|
||||
"attributableFailureCount": 12,
|
||||
"attributableFailureCount": 13,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 12,
|
||||
"decisionTotal": 13,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 166,
|
||||
"framework_architecture_unsupported": 11,
|
||||
"framework_architecture_unsupported": 12,
|
||||
"memory_capacity": 1,
|
||||
"参数/模板问题": 18,
|
||||
"日志缺失": 1,
|
||||
"验证失败": 23
|
||||
},
|
||||
"failureCount": 220,
|
||||
"failureCount": 221,
|
||||
"failureRate": 1.0,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"total": 220,
|
||||
"total": 221,
|
||||
"unresolvedFailureCount": 208
|
||||
},
|
||||
"Kunlunxin_r-200-8f": {
|
||||
@@ -5880,7 +5901,7 @@
|
||||
"hygon_k100-ai": 64.0
|
||||
},
|
||||
"pendingRecords": 0,
|
||||
"policyCancelledRecords": 182,
|
||||
"policyCancelledRecords": 183,
|
||||
"profileCombinationStats": {
|
||||
"Ascend_910-b3|vllm_tokenizer_patch|text-generation|bailing_hybrid|none": {
|
||||
"attributableFailureCount": 1,
|
||||
@@ -13015,6 +13036,29 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation|spark2_5|none": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"modelType": "spark2_5",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Kunlunxin_p-800",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Kunlunxin_p-800|vllm|text-generation|deepseek_v4|none": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
@@ -15307,10 +15351,9 @@
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1,
|
||||
"platform_infrastructure": 1
|
||||
},
|
||||
"failureCount": 2,
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"lastPlatformFailureAt": "2026-09-20T03:41:42.246689+00:00",
|
||||
@@ -15322,8 +15365,8 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b4",
|
||||
"taskType": "text-generation",
|
||||
"total": 2,
|
||||
"unresolvedFailureCount": 1
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Ascend_910-b4|vllm|text-generation": {
|
||||
"attributableFailureCount": 2,
|
||||
@@ -15773,17 +15816,17 @@
|
||||
"unresolvedFailureCount": 2
|
||||
},
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation": {
|
||||
"attributableFailureCount": 4,
|
||||
"consecutiveFailures": 4,
|
||||
"attributableFailureCount": 5,
|
||||
"consecutiveFailures": 5,
|
||||
"consecutivePlatformFailures": 0,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 4,
|
||||
"decisionTotal": 5,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 4,
|
||||
"framework_architecture_unsupported": 5,
|
||||
"参数/模板问题": 10
|
||||
},
|
||||
"failureCount": 14,
|
||||
"failureCount": 15,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"lastPlatformFailureAt": null,
|
||||
@@ -15795,7 +15838,7 @@
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Kunlunxin_p-800",
|
||||
"taskType": "text-generation",
|
||||
"total": 14,
|
||||
"total": 15,
|
||||
"unresolvedFailureCount": 10
|
||||
},
|
||||
"MetaX_c-500|unknown|text-generation": {
|
||||
@@ -16283,31 +16326,6 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|phi3|compressed-tensors": {
|
||||
"attributableFailureCount": 0,
|
||||
"consecutiveFailures": 0,
|
||||
"decisionFailureRate": 0.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 0,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"lastTerminalAt": "2026-09-21T00:42:10.765015+00:00",
|
||||
"modelType": "phi3",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "compressed-tensors",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Ascend_910-b4",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 1
|
||||
},
|
||||
"Ascend_910-b4|vllm_tokenizer_patch|text-generation|qwen2|compressed-tensors": {
|
||||
"attributableFailureCount": 0,
|
||||
"consecutiveFailures": 0,
|
||||
@@ -19314,6 +19332,31 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation|spark2_5|none": {
|
||||
"attributableFailureCount": 1,
|
||||
"consecutiveFailures": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"lastTerminalAt": "2026-09-21T08:25:31.485651+00:00",
|
||||
"modelType": "spark2_5",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Kunlunxin_p-800",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"MetaX_c-500|vllm|text-generation|gemma3|compressed-tensors": {
|
||||
"attributableFailureCount": 1,
|
||||
"consecutiveFailures": 1,
|
||||
@@ -29780,6 +29823,30 @@
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Kunlunxin_p-800|vllm_tokenizer_patch|text-generation|spark2_5|none|31": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
"decisionSuccessRate": 0.0,
|
||||
"decisionTotal": 1,
|
||||
"failureBreakdown": {
|
||||
"framework_architecture_unsupported": 1
|
||||
},
|
||||
"failureCount": 1,
|
||||
"failureRate": 1.0,
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"loadSizeLog2Bucket": 31,
|
||||
"modelType": "spark2_5",
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 0,
|
||||
"quantizationMethod": "none",
|
||||
"successCount": 0,
|
||||
"successRate": 0.0,
|
||||
"targetGpu": "Kunlunxin_p-800",
|
||||
"taskType": "text-generation",
|
||||
"total": 1,
|
||||
"unresolvedFailureCount": 0
|
||||
},
|
||||
"Kunlunxin_p-800|vllm|text-generation|deepseek_v4|none|36": {
|
||||
"attributableFailureCount": 1,
|
||||
"decisionFailureRate": 1.0,
|
||||
@@ -33012,20 +33079,20 @@
|
||||
"unresolvedFailureCount": 0
|
||||
}
|
||||
},
|
||||
"terminalRecords": 16185,
|
||||
"totalRecords": 16367,
|
||||
"terminalRecords": 16186,
|
||||
"totalRecords": 16369,
|
||||
"totals": {
|
||||
"attributableFailureCount": 5861,
|
||||
"decisionFailureRate": 0.8615,
|
||||
"decisionSuccessRate": 0.1385,
|
||||
"decisionTotal": 6803,
|
||||
"attributableFailureCount": 5862,
|
||||
"decisionFailureRate": 0.8616,
|
||||
"decisionSuccessRate": 0.1384,
|
||||
"decisionTotal": 6804,
|
||||
"failureBreakdown": {
|
||||
"ambiguous_runtime": 3894,
|
||||
"architecture_compatibility": 212,
|
||||
"attention_backend": 1,
|
||||
"backend_operator": 102,
|
||||
"context_length": 319,
|
||||
"framework_architecture_unsupported": 2056,
|
||||
"framework_architecture_unsupported": 2057,
|
||||
"memory_capacity": 1197,
|
||||
"model_load": 495,
|
||||
"platform_infrastructure": 924,
|
||||
@@ -33036,14 +33103,14 @@
|
||||
"日志缺失": 719,
|
||||
"验证失败": 674
|
||||
},
|
||||
"failureCount": 15243,
|
||||
"failureCount": 15244,
|
||||
"failureRate": 0.9418,
|
||||
"pendingCount": 0,
|
||||
"pendingRate": 0.0,
|
||||
"platformFailureCount": 924,
|
||||
"successCount": 942,
|
||||
"successRate": 0.0582,
|
||||
"total": 16185,
|
||||
"total": 16186,
|
||||
"unresolvedFailureCount": 8458
|
||||
},
|
||||
"warnings": [
|
||||
@@ -33100,6 +33167,6 @@
|
||||
]
|
||||
},
|
||||
"storageMode": "decision_state_only",
|
||||
"summarizedRecords": 16367,
|
||||
"summarizedRecords": 16369,
|
||||
"version": 1
|
||||
}
|
||||
|
||||
@@ -20,7 +20,7 @@
|
||||
"cleanupDisabled": true,
|
||||
"reason": "admission_only"
|
||||
},
|
||||
"architectureBlockCount": 132,
|
||||
"architectureBlockCount": 133,
|
||||
"architectureFrameworkCatalog": {
|
||||
"ascend_910-b3|text-generation": [
|
||||
"llamacpp",
|
||||
@@ -73,8 +73,30 @@
|
||||
]
|
||||
},
|
||||
"architectureFrameworkCatalogErrors": {},
|
||||
"architectureIncompatibleCount": 0,
|
||||
"architectureIncompatibleTasks": [],
|
||||
"architectureIncompatibleCount": 1,
|
||||
"architectureIncompatibleTasks": [
|
||||
{
|
||||
"accountIndex": 5,
|
||||
"architectureBlockEvidenceCount": 1,
|
||||
"architectureBlockExpiresAt": "2026-10-20T03:25:31.991647+00:00",
|
||||
"architectureMatchScope": "exact_framework",
|
||||
"architectureMatchType": "architectures",
|
||||
"architectureSignature": "architectures:spark2_5forcausallm",
|
||||
"architectureSignatures": [
|
||||
"architectures:spark2_5forcausallm"
|
||||
],
|
||||
"evaluatedFrameworks": [
|
||||
"vllm_tokenizer_patch"
|
||||
],
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"gpuType": "Kunlunxin_p-800",
|
||||
"modelId": "XHToken/Spark-X2.5-4B-FP8",
|
||||
"reason": "known_framework_architecture_incompatible",
|
||||
"status": "waiting",
|
||||
"taskId": 4969629,
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
],
|
||||
"architectureModelConfigErrors": {
|
||||
"GestaltLabs/Ornstein-3.5-9B-V2-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/GestaltLabs/Ornstein-3.5-9B-V2-GGUF/resolve/master/config.json (status=404)",
|
||||
"LiquidAI/LFM2.5-2.6B-DSpark-GGUF": "HttpJsonError: HTTP request failed for https://modelscope.cn/models/LiquidAI/LFM2.5-2.6B-DSpark-GGUF/resolve/master/config.json (status=500)",
|
||||
@@ -103,17 +125,42 @@
|
||||
"frameworkCatalogUnknown": 0,
|
||||
"frameworkContextUnknown": 66,
|
||||
"modelArchitectureUnknown": 133,
|
||||
"noMatchingBlock": 990,
|
||||
"noMatchingBlock": 989,
|
||||
"partiallyBlockedFrameworkSet": 21,
|
||||
"runningMatchedProtected": 0,
|
||||
"submissionContextMismatch": 0,
|
||||
"submissionContextUnknown": 0
|
||||
},
|
||||
"cancelledCount": 0,
|
||||
"cancelledTasks": [],
|
||||
"cancelledCount": 1,
|
||||
"cancelledTasks": [
|
||||
{
|
||||
"accountIndex": 5,
|
||||
"architectureBlockEvidenceCount": 1,
|
||||
"architectureBlockExpiresAt": "2026-10-20T03:25:31.991647+00:00",
|
||||
"architectureMatchScope": "exact_framework",
|
||||
"architectureMatchType": "architectures",
|
||||
"architectureSignature": "architectures:spark2_5forcausallm",
|
||||
"architectureSignatures": [
|
||||
"architectures:spark2_5forcausallm"
|
||||
],
|
||||
"cleanupReasons": [
|
||||
"known_framework_architecture_incompatible"
|
||||
],
|
||||
"evaluatedFrameworks": [
|
||||
"vllm_tokenizer_patch"
|
||||
],
|
||||
"framework": "vllm_tokenizer_patch",
|
||||
"gpuType": "Kunlunxin_p-800",
|
||||
"modelId": "XHToken/Spark-X2.5-4B-FP8",
|
||||
"reason": "known_framework_architecture_incompatible",
|
||||
"status": "waiting",
|
||||
"taskId": 4969629,
|
||||
"taskType": "text-generation"
|
||||
}
|
||||
],
|
||||
"certainOomCount": 0,
|
||||
"certainOomTasks": [],
|
||||
"cleanupCandidateCount": 0,
|
||||
"cleanupCandidateCount": 1,
|
||||
"dryRun": false,
|
||||
"listingErrors": {},
|
||||
"modelAgeErrors": {},
|
||||
@@ -138,7 +185,7 @@
|
||||
],
|
||||
"oldOverflowCount": 0,
|
||||
"oldOverflowTasks": [],
|
||||
"policyCancelledRecorded": 0,
|
||||
"policyCancelledRecorded": 1,
|
||||
"policyNoLongerAppliesCount": 0,
|
||||
"policyNoLongerAppliesTasks": [],
|
||||
"recentModelDays": 7,
|
||||
|
||||
@@ -250,6 +250,7 @@
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm", "lastSyncTime": "2026-09-20T17:09:45.266076+00:00", "modelId": "neuralmagic/Mistral-7B-Instruct-v0.3-quantized.w8a8", "modelProfile": {"architectures": ["MistralForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7519537296, "estimatedRequiredGiB": 8.407, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral", "modelscopeFileSize": 7522510264, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7248023552, "modelscopeTags": ["license:apache-2.0", "model_type:mistral", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7522510264}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:25:32.254780+00:00", "targetGpu": "Biren_166m", "taskId": "4969646", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "参数/模板问题", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T01:08:25.664296+00:00", "modelId": "Youssofal/Qwen3.5-9B-MTPLX-Optimized-Speed", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T03:25:32.245630+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969647", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": "参数/模板问题", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T02:15:17.867315+00:00", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Speed-FP16", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T03:25:31.995522+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969637", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedArchitectures": ["Spark2_5ForCausalLM"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T08:25:31.485651+00:00", "modelId": "XHToken/Spark-X2.5-1.7B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3415340496, "estimatedRequiredGiB": 3.834, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 3430847031, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1707657216, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3430847031}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:25:31.991647+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969636", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T00:42:10.765050+00:00", "modelId": "JANGQ-AI/AppleScript-8B-JANG_4M", "modelProfile": {"architectures": ["ZayaForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7915599888, "estimatedRequiredGiB": 8.885, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "zaya", "modelscopeFileSize": 7950463599, "modelscopeLicense": "other", "modelscopeParams": 2274126328, "modelscopeTags": ["license:other", "model_type:zaya", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:applescript", "custom_tag:macos", "custom_tag:computer-use", "custom_tag:agent", "custom_tag:tool-calling", "custom_tag:function-calling", "custom_tag:mlx", "custom_tag:moe", "custom_tag:osaurus"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7950463599}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:25:24.240446+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969618", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "execute_empty_result", "failureCode": "EXECUTE_EMPTY_RESULT", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T00:42:10.765072+00:00", "modelId": "OpenBMB/BitCPM-CANN-8B", "modelProfile": {"architectures": ["MiniCPMForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16370604914, "estimatedRequiredGiB": 18.305, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "minicpm", "modelscopeFileSize": 16378634956, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:minicpm", "library:pytorch", "library:transformer", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16378634956}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:25:24.238793+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969621", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "参数/模板问题", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T01:08:25.664243+00:00", "modelId": "logic65/Qwen3.8-Whittle-w50-18.3B-unrepaired", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T03:24:23.636965+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969578", "taskType": "text-generation", "verifyResult": null}
|
||||
@@ -297,4 +298,3 @@
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_text"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T03:45:50.542829+00:00", "modelId": "TokenRhythm/NeoHorse-1-4B", "modelProfile": {"architectures": ["Qwen3_5ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8411552560, "estimatedRequiredGiB": 9.428, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_text", "modelscopeFileSize": 8435794426, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 4205751296, "modelscopeTags": ["license:Apache License 2.0", "model_type:qwen3_5_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agentic", "custom_tag:tool-use", "custom_tag:coding", "custom_tag:reasoning", "custom_tag:instruction-following"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8435794426}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:17.849368+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969324", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "framework_architecture_unsupported", "failureAction": "block_gpu_framework_architecture", "failureCategory": "framework_architecture_unsupported", "failureClassificationReason": "explicit_framework_model_unsupported", "failureCode": "MODEL_NOT_SUPPORTED", "failureDetectedFramework": "vllm", "failureDeterministic": true, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": false, "failureScope": "model_gpu_framework", "failureUnsupportedModelTypes": ["qwen3_5_moe"], "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T20:49:19.683330+00:00", "modelId": "apodex/Apodex-1.1-mini-GPTQ-Int4", "modelProfile": {"architectures": ["Qwen3_5MoeForConditionalGeneration"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 24651300904, "estimatedRequiredGiB": 27.587, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe", "modelscopeFileSize": 24684544516, "modelscopeLicense": "apache-2.0", "modelscopeParams": 35951822704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:apodex", "custom_tag:arxiv:2608.23283"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "gptq", "repositoryOnDiskBytes": 24684544516}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:17.835766+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969320", "taskType": "text-generation", "verifyResult": -1}
|
||||
{"failReason": "参数/模板问题", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T02:15:17.867242+00:00", "modelId": "ornith-ai/Ornith-1.5-9B-NVFP4", "modelProfile": {}, "outcome": "failed", "status": "failed", "submitTime": "2026-09-20T03:09:17.795054+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969322", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": "ambiguous_runtime", "failureAction": "send_compact_profile_and_root_exception_to_llm", "failureCategory": "ambiguous_runtime", "failureClassificationReason": "driver_error", "failureCode": "DRIVER_ERROR", "failureDetectedFramework": "vllm", "failureDeterministic": false, "failureEnrichmentAttempts": 1, "failureEnrichmentError": null, "failureNeedsLlm": true, "failureScope": "unknown", "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-21T00:42:10.765015+00:00", "modelId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w4a16", "modelProfile": {"architectures": ["Phi3ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7686303632, "estimatedRequiredGiB": 8.592, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "phi3", "modelscopeFileSize": 7688299781, "modelscopeLicense": "llama2", "modelscopeParams": 13960238080, "modelscopeTags": ["license:llama2", "model_type:phi3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7688299781}, "outcome": "failed", "status": "success", "submitTime": "2026-09-20T03:09:11.542224+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4969301", "taskType": "text-generation", "verifyResult": -1}
|
||||
|
||||
@@ -2003,6 +2003,63 @@
|
||||
{"batchId": "94fa65280ebc4ffe87b819cfb4433489", "completedAt": "2026-09-21T08:17:20.360601+00:00", "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:17:10.101584+00:00", "framework": "vllm_tokenizer_patch", "intentId": "0ae64386f6ca48c38accae5cd676d52c", "lastModified": "2026-09-17T15:15:04+00:00", "modelAddress": "https://modelscope.cn/models/mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "reason": null, "reconciledAt": "2026-09-21T08:22:13.016425+00:00", "repoId": "mlx-community/gemma-4-31B-it-qat-OptiQ-4bit", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Kunlunxin_p-800", "taskId": "4996859", "taskType": "text-generation"}
|
||||
{"batchId": "94fa65280ebc4ffe87b819cfb4433489", "completedAt": "2026-09-21T08:17:20.360620+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:17:10.101734+00:00", "framework": "vllm", "intentId": "14a017ca764f4d31a47c1b218150d2fb", "lastModified": "2026-09-17T13:06:10+00:00", "modelAddress": "https://modelscope.cn/models/Youssofal/Qwen3.8-27B-MTPLX-Optimized-Quality", "reason": null, "reconciledAt": "2026-09-21T08:22:13.016153+00:00", "repoId": "Youssofal/Qwen3.8-27B-MTPLX-Optimized-Quality", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "recovered_active", "targetGpu": "Kunlunxin_p-800", "taskId": "4996858", "taskType": "text-generation"}
|
||||
{"batchId": "94fa65280ebc4ffe87b819cfb4433489", "completedAt": "2026-09-21T08:17:20.360624+00:00", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:17:10.101799+00:00", "framework": "vllm_fix_tokenizer", "intentId": "1819d015d79041a98b8592001f18c384", "lastModified": "2026-09-21T01:31:41+00:00", "modelAddress": "https://modelscope.cn/models/blue2star/Qwen-Image-2.1-PE-T2I-ComfyUI", "reason": null, "reconciledAt": "2026-09-21T08:22:13.019512+00:00", "repoId": "blue2star/Qwen-Image-2.1-PE-T2I-ComfyUI", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Biren_166m", "taskId": "4996857", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341582+00:00", "framework": "vllm_tokenizer_patch", "intentId": "4ecceef8c3e34e69b5633c14fd76e0d6", "lastModified": "2026-08-24T18:24:11+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/tinyllama-oneshot-w8a8-dynamic-token-v2", "repoId": "nm-testing/tinyllama-oneshot-w8a8-dynamic-token-v2", "safeConfigVector": {"gpuNum": 1}, "status": "pending", "targetGpu": "Ascend_910-b4", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341718+00:00", "framework": "vllm_tokenizer_patch", "intentId": "35a96e7094bc4b2ab305fa4795c8326e", "lastModified": "2026-08-24T19:10:38+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/SmolLM-135M-Instruct-quantized.w8a16", "repoId": "RedHatAI/SmolLM-135M-Instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1}, "status": "pending", "targetGpu": "Ascend_910-b4", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "c25be7d7c1bf173abe564916f8a370d9b648beb3aea2a79522523aaa6b69ee29", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341772+00:00", "framework": "vllm", "intentId": "285e69eaf8da4d1fa03807805fd02a33", "lastModified": "2026-08-24T18:26:39+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/tinyllama-oneshot-w8a8-static-v2", "repoId": "nm-testing/tinyllama-oneshot-w8a8-static-v2", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "c25be7d7c1bf173abe564916f8a370d9b648beb3aea2a79522523aaa6b69ee29", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341822+00:00", "framework": "vllm", "intentId": "d68a85d02d64402f97de86717c422232", "lastModified": "2026-08-24T18:28:40+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/tinyllama-marlin24-w4a16-group128", "repoId": "nm-testing/tinyllama-marlin24-w4a16-group128", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341871+00:00", "framework": "vllm", "intentId": "aef435d05a004cc495a1e661cc56a4f9", "lastModified": "2026-08-26T16:45:23+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Gemma-SEA-LION-v4-27B-IT-NVFP4", "repoId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341919+00:00", "framework": "vllm_tokenizer_patch", "intentId": "d7fe2a4d519a43b7a87e57f06ed6c672", "lastModified": "2026-08-24T19:39:27+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-3b-FP8", "repoId": "neuralmagic/starcoder2-3b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "26bfee231695e882b96dee0191bc1dfec25bcad132af96a7ee5f0e26f3cb9fdc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341960+00:00", "framework": "vllm_tokenizer_patch", "intentId": "cee162f9da1e4e5c95e343a65fc35a18", "lastModified": "2026-08-24T18:09:20+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/tinyllama-oneshot-w8w8-test-static-shape-change", "repoId": "nm-testing/tinyllama-oneshot-w8w8-test-static-shape-change", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.341999+00:00", "framework": "vllm_tokenizer_patch", "intentId": "cd499973f8c14a04bb930a55709b7b05", "lastModified": "2026-08-24T18:15:08+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W8-Channel-A8-Dynamic-Per-Token-Test", "repoId": "nm-testing/Meta-Llama-3-8B-Instruct-W8-Channel-A8-Dynamic-Per-Token-Test", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342037+00:00", "framework": "vllm_tokenizer_patch", "intentId": "ec983485c729450196906ea0bd55c7b9", "lastModified": "2026-08-24T18:12:10+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/Meta-Llama-3-8B-Instruct-W8A8-FP8-Channelwise-compressed-tensors", "repoId": "nm-testing/Meta-Llama-3-8B-Instruct-W8A8-FP8-Channelwise-compressed-tensors", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "26bfee231695e882b96dee0191bc1dfec25bcad132af96a7ee5f0e26f3cb9fdc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342074+00:00", "framework": "vllm_tokenizer_patch", "intentId": "b5c8312512a44121b2bab7b1eb657a99", "lastModified": "2026-08-24T18:26:53+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/tinyllama-oneshot-w8a16-per-channel", "repoId": "nm-testing/tinyllama-oneshot-w8a16-per-channel", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342112+00:00", "framework": "vllm_tokenizer_patch", "intentId": "d6f012237fc44b4ca50964d7564ce9c2", "lastModified": "2026-08-24T18:26:59+00:00", "modelAddress": "https://modelscope.cn/models/nm-testing/nonuniform", "repoId": "nm-testing/nonuniform", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342151+00:00", "framework": "vllm_tokenizer_patch", "intentId": "c325c22a7b5243718e5b4802060b4a78", "lastModified": "2026-08-24T19:41:12+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-15b-quantized.w8a8", "repoId": "RedHatAI/starcoder2-15b-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342188+00:00", "framework": "vllm_tokenizer_patch", "intentId": "a865966dac654e718af93b88021a317d", "lastModified": "2026-08-24T19:34:42+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-FP8", "repoId": "RedHatAI/Phi-3-mini-128k-instruct-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342226+00:00", "framework": "vllm_tokenizer_patch", "intentId": "119fadf5a18f49219be7591c4ab05efe", "lastModified": "2026-08-24T20:05:18+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "repoId": "RedHatAI/Meta-Llama-3.1-8B-Instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342264+00:00", "framework": "vllm_tokenizer_patch", "intentId": "1e2f4ae26340433ba848c828a32c8037", "lastModified": "2026-08-24T19:59:13+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-FP8", "repoId": "RedHatAI/starcoder2-7b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342301+00:00", "framework": "vllm_tokenizer_patch", "intentId": "8a9d6c4b42ec4fbabc1451982b668aa0", "lastModified": "2026-08-24T19:38:03+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "repoId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342339+00:00", "framework": "vllm_tokenizer_patch", "intentId": "fd21d9669f0344bf92bfba43f2718fea", "lastModified": "2026-08-24T19:51:52+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "repoId": "RedHatAI/Phi-3-small-128k-instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342376+00:00", "framework": "vllm_tokenizer_patch", "intentId": "e4bd52679b1548b983ed291e9a6bd3d7", "lastModified": "2026-08-24T19:08:49+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w8a16", "repoId": "RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342414+00:00", "framework": "vllm_tokenizer_patch", "intentId": "9a25b374aa534e66b5814c8cbc3d9944", "lastModified": "2026-08-24T19:36:07+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w4a16", "repoId": "RedHatAI/Mistral-Nemo-Instruct-2407-quantized.w4a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "26bfee231695e882b96dee0191bc1dfec25bcad132af96a7ee5f0e26f3cb9fdc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342451+00:00", "framework": "vllm_tokenizer_patch", "intentId": "c19c291c800b4cab8cd93cae146073cd", "lastModified": "2026-08-24T19:35:49+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/SmolLM-1.7B-Instruct-quantized.w8a16", "repoId": "RedHatAI/SmolLM-1.7B-Instruct-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342489+00:00", "framework": "vllm_tokenizer_patch", "intentId": "b63bc3a1f1d24842bb163e1ffdbf2780", "lastModified": "2026-08-24T19:06:55+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-2b-it-quantized.w8a8", "repoId": "RedHatAI/gemma-2-2b-it-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342527+00:00", "framework": "vllm_tokenizer_patch", "intentId": "1155224ce24c4c9588168a7dadcdd978", "lastModified": "2026-08-24T20:02:01+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-2b-quantized.w8a16", "repoId": "RedHatAI/gemma-2-2b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Ascend_910-b3", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342565+00:00", "framework": "vllm", "intentId": "bde149fbc27343698ea399f18b1d5292", "lastModified": "2026-08-24T22:55:11+00:00", "modelAddress": "https://modelscope.cn/models/BAAI/AquilaMed-RL", "repoId": "BAAI/AquilaMed-RL", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342625+00:00", "framework": "vllm", "intentId": "28f5efe1a3824e05ba64c037f54e78df", "lastModified": "2026-08-24T19:26:17+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w4a16", "repoId": "neuralmagic/gemma-2-2b-it-quantized.w4a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342688+00:00", "framework": "vllm", "intentId": "9a32c56e1ccd427f9b5449323760dcd5", "lastModified": "2026-08-25T12:06:29+00:00", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-8B-A1B", "repoId": "LiquidAI/LFM2.5-8B-A1B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342745+00:00", "framework": "vllm", "intentId": "8a389d95d3294278ad2fc527e21619ec", "lastModified": "2026-08-26T20:08:57+00:00", "modelAddress": "https://modelscope.cn/models/LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "repoId": "LiquidAI/LFM2.5-1.2B-Thinking-MLX-6bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342802+00:00", "framework": "vllm", "intentId": "acd79326fcd44084bc5b9a5f04fa9602", "lastModified": "2026-08-24T19:25:45+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-2-7b-chat-quantized.w8a8", "repoId": "RedHatAI/Llama-2-7b-chat-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342859+00:00", "framework": "vllm", "intentId": "2872bffdedf5407e8ca379ee586510d2", "lastModified": "2026-08-24T20:17:00+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "repoId": "RedHatAI/Phi-3-mini-128k-instruct-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342916+00:00", "framework": "vllm", "intentId": "76ed48a71a5c42849406786189b71738", "lastModified": "2026-08-24T20:08:44+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "repoId": "RedHatAI/Phi-3-medium-128k-instruct-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.342972+00:00", "framework": "vllm", "intentId": "ba014286cc784b99b5a7a8eaaf4c1a58", "lastModified": "2026-09-10T07:49:50+00:00", "modelAddress": "https://modelscope.cn/models/TokenRhythm/NeoHorse-1-9B", "repoId": "TokenRhythm/NeoHorse-1-9B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343029+00:00", "framework": "vllm", "intentId": "be4c50873a504559badea171353a2aff", "lastModified": "2026-09-08T13:43:47+00:00", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/Osaurus-AppleScript-8B-JANG_6M", "repoId": "JANGQ-AI/Osaurus-AppleScript-8B-JANG_6M", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343086+00:00", "framework": "vllm", "intentId": "733b3874eef84ec891d1efc803f20359", "lastModified": "2026-08-26T16:41:36+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "repoId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343142+00:00", "framework": "vllm", "intentId": "a19a92ee848e408e9a96946470b5428c", "lastModified": "2026-09-08T15:40:32+00:00", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/Ling-2.6-flash-JANGTQ", "repoId": "JANGQ-AI/Ling-2.6-flash-JANGTQ", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343200+00:00", "framework": "vllm", "intentId": "a438a711d43f492e81dd6b2f21448e6b", "lastModified": "2026-08-26T15:57:06+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "repoId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343256+00:00", "framework": "vllm", "intentId": "c56ab15797944ffc876085d5c001e0ec", "lastModified": "2026-09-09T06:27:24+00:00", "modelAddress": "https://modelscope.cn/models/ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "repoId": "ddalcu/Muse-Glimmer-30B-MLX-Serve-4bit", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343314+00:00", "framework": "vllm", "intentId": "009b98968a364eb89d04269af6378068", "lastModified": "2026-08-26T18:16:42+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/SEA-LION-v1-7B-IT-Research", "repoId": "aisingapore/SEA-LION-v1-7B-IT-Research", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "pending", "targetGpu": "Kunlunxin_p-800", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343371+00:00", "framework": "vllm_fix_tokenizer", "intentId": "4163cad4e0fa40d5898ae5199def7ac3", "lastModified": "2026-08-31T03:35:12+00:00", "modelAddress": "https://modelscope.cn/models/logic65/Qwen3.8-Whittle-w50-18.3B-unrepaired", "repoId": "logic65/Qwen3.8-Whittle-w50-18.3B-unrepaired", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343416+00:00", "framework": "vllm_fix_tokenizer", "intentId": "20cdbe73f7054e7699a7e97c2ef51820", "lastModified": "2026-09-10T14:58:25+00:00", "modelAddress": "https://modelscope.cn/models/webAI-Official/TwIL-LM3", "repoId": "webAI-Official/TwIL-LM3", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Biren_166m", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343459+00:00", "framework": "vllm", "intentId": "1a933c3a0ee047c193356e652c151e56", "lastModified": "2026-08-24T19:21:33+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-3b-quantized.w8a8", "repoId": "neuralmagic/starcoder2-3b-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Iluvatar_mrv-100", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343508+00:00", "framework": "vllm-patch-tokenizer", "intentId": "7f90348806264016a67ae33ddc5346de", "lastModified": "2026-08-24T19:53:06+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "repoId": "neuralmagic/Phi-3-medium-128k-instruct-quantized.w4a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343551+00:00", "framework": "vllm-patch-tokenizer", "intentId": "59aea16fe95a404287465515a9cde451", "lastModified": "2026-08-24T20:09:59+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Llama-2-7b-chat-quantized.w8a8", "repoId": "neuralmagic/Llama-2-7b-chat-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343595+00:00", "framework": "vllm-patch-tokenizer", "intentId": "e257e711259c4b21ab73fb9634c853ba", "lastModified": "2026-08-24T20:10:58+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Meta-Llama-3.1-8B-quantized.w8a16", "repoId": "neuralmagic/Meta-Llama-3.1-8B-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343637+00:00", "framework": "vllm-patch-tokenizer", "intentId": "681273e1ad22435fa576d4c6dcc9e3e0", "lastModified": "2026-08-24T19:42:52+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-7b-FP8", "repoId": "neuralmagic/starcoder2-7b-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343689+00:00", "framework": "vllm-patch-tokenizer", "intentId": "5a56d2657e0d40db934210336d750a4a", "lastModified": "2026-08-24T20:21:52+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-7b-quantized.w8a16", "repoId": "neuralmagic/starcoder2-7b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343731+00:00", "framework": "vllm-patch-tokenizer", "intentId": "9037e7ff7e9b4e05abf4ad7401dfbaaa", "lastModified": "2026-08-24T19:39:20+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-3b-quantized.w8a16", "repoId": "neuralmagic/starcoder2-3b-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343774+00:00", "framework": "vllm-patch-tokenizer", "intentId": "b932617a80f9453a888ce97be73f89c4", "lastModified": "2026-08-24T20:09:49+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/Llama-3.2-3B-Instruct-FP8", "repoId": "RedHatAI/Llama-3.2-3B-Instruct-FP8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343815+00:00", "framework": "vllm-patch-tokenizer", "intentId": "23d4dc98d3f34c06bcab5a1d43c7e0cc", "lastModified": "2026-08-24T19:20:59+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-9b-it-quantized.w8a8", "repoId": "RedHatAI/gemma-2-9b-it-quantized.w8a8", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343858+00:00", "framework": "vllm-patch-tokenizer", "intentId": "36a7c5f66b9e42e1bee3d5d4ebb8aac5", "lastModified": "2026-08-24T20:18:47+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/gemma-2-27b-it-quantized.w8a16", "repoId": "RedHatAI/gemma-2-27b-it-quantized.w8a16", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "hygon_k100-ai", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343900+00:00", "framework": "vllm", "intentId": "a55b2474f6e24cc99d53517cfc430255", "lastModified": "2026-08-24T20:10:55+00:00", "modelAddress": "https://modelscope.cn/models/RedHatAI/starcoder2-7b-quantized.w8a8", "repoId": "RedHatAI/starcoder2-7b-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "ef4c08c913d3f64cae22550f0c4a189795901d4b11e26d8e640e07e857659fb9", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343949+00:00", "framework": "vllm", "intentId": "0a3522ffa8c04cbb9e4ac978a1138911", "lastModified": "2026-09-08T15:45:26+00:00", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/AppleScript-8B-JANG_4M", "repoId": "JANGQ-AI/AppleScript-8B-JANG_4M", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "f49713b9fe8e3e7683ab2722544c85a44993bcd72b9421c1599f634e68b3a2f7", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.343996+00:00", "framework": "vllm", "intentId": "9cbf4d8fb2df4c1eb256ac2cc6535a25", "lastModified": "2026-08-26T17:29:18+00:00", "modelAddress": "https://modelscope.cn/models/aisingapore/SEA-LION-v1-7B", "repoId": "aisingapore/SEA-LION-v1-7B", "safeConfigVector": {"dtype": "-", "gpuNum": 1}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x4", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.344043+00:00", "framework": "vllm-customized", "intentId": "04f06cc5707a4f009b8f8a0e1550e4b7", "lastModified": "2026-09-09T12:19:20+00:00", "modelAddress": "https://modelscope.cn/models/OpenBMB/BitCPM-CANN-8B", "repoId": "OpenBMB/BitCPM-CANN-8B", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.344103+00:00", "framework": "vllm-customized", "intentId": "17777273175d453fb4b3763ecc109c4f", "lastModified": "2026-08-24T20:20:16+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/starcoder2-15b-FP8", "repoId": "neuralmagic/starcoder2-15b-FP8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.344161+00:00", "framework": "vllm-customized", "intentId": "aa7b8f907336466a87d56ebcb5548e69", "lastModified": "2026-08-24T20:01:20+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/Llama-3.2-1B-Instruct-quantized.w8a8", "repoId": "neuralmagic/Llama-3.2-1B-Instruct-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.344219+00:00", "framework": "vllm-customized", "intentId": "3c985f6f5d6a41dba83c9b4291b30e6f", "lastModified": "2026-08-24T20:22:00+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w8a16", "repoId": "neuralmagic/gemma-2-2b-it-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.344276+00:00", "framework": "vllm-customized", "intentId": "5a5dbc3f15524ec5a1078de10692df88", "lastModified": "2026-08-26T17:18:42+00:00", "modelAddress": "https://modelscope.cn/models/nv-community/NVIDIA-Nemotron-3-Nano-30B-A3B-NVFP4", "repoId": "nv-community/NVIDIA-Nemotron-3-Nano-30B-A3B-NVFP4", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||
{"batchId": "53c59e9635fa4b4abb2faa302ec48e2b", "configFingerprint": "1b41437401598cde7836ccc25883635c2738b06a5af924b7497b6e747e655b4b", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:29:14.344333+00:00", "framework": "vllm-customized", "intentId": "6361228927444e939da28960c5f8f929", "lastModified": "2026-08-24T19:44:39+00:00", "modelAddress": "https://modelscope.cn/models/neuralmagic/gemma-2-2b-it-quantized.w8a8", "repoId": "neuralmagic/gemma-2-2b-it-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 4096}, "status": "pending", "targetGpu": "Cambricon_mlu-370-x8", "taskType": "text-generation"}
|
||||
{"batchId": "7dd2dfc02db643c1b324bbebd0018229", "completedAt": "2026-09-21T08:16:22.466082+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:11:15.127288+00:00", "framework": "vllm", "intentId": "65cef4d446084037bb13153d52b7c872", "repoId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-4BIT", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "age_policy_deferred", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
|
||||
{"batchId": "7dd2dfc02db643c1b324bbebd0018229", "completedAt": "2026-09-21T08:16:22.466079+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:11:15.127231+00:00", "framework": "vllm", "intentId": "a7743eff51694235ba4577c5edb0a5c8", "repoId": "RedHatAI/gemma-2-2b-it-quantized.w8a16", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "age_policy_deferred", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
|
||||
{"batchId": "7dd2dfc02db643c1b324bbebd0018229", "completedAt": "2026-09-21T08:16:22.466076+00:00", "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configSource": "modelhub_live", "createdAt": "2026-09-21T08:11:15.127174+00:00", "framework": "vllm", "intentId": "d8aff6a5552345c6aabfa8a2186a5135", "repoId": "RedHatAI/Qwen2-1.5B-Instruct-quantized.w8a8", "safeConfigVector": {"dtype": "-", "gpuNum": 1, "maxModelLen": 2048}, "status": "age_policy_deferred", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"}
|
||||
|
||||
@@ -2,24 +2,24 @@
|
||||
"agentVersion": "2026.09.20.2",
|
||||
"checksums": {
|
||||
".modelhub_state/account_capacity.json": "68d2a8390ad022f2ddcd30841f3860c4535387fa64cb46dc7507eb16837e798f",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "c9b60327e144fd317198bb237e675a6bd83df744fb84dfd813a38b38f2deb556",
|
||||
".modelhub_state/architecture_compatibility_blacklist.json": "e3137b051a76a893cd9c323d311d8aec5731cfc13d593eb63753065334b00e99",
|
||||
".modelhub_state/architecture_history_backfill.json": "a92348206605da04f75c9e26e7c818b76f259e0b28983f18b6375ef94802f0ab",
|
||||
".modelhub_state/market_intelligence.json": "43f9d24f7580107a2db5df7ea75ae8b87627b482c8e896d7f4fc9417d594ce86",
|
||||
".modelhub_state/official_capabilities.json": "ccdeb985262abdd9f3daea044d8baac85ca515700520d5b8c515f5e8615b1c08",
|
||||
".modelhub_state/outcome_checkpoint.json": "208c35feca4ca68606616c39d95138be80ed3c269d30591a7237adc2c4fb3068",
|
||||
".modelhub_state/queue_cleanup_latest.json": "f55f6627ea14e95953fea4643ff04c475523c53a2ebada48870de5a460ae306e",
|
||||
".modelhub_state/recent_outcomes.jsonl": "46878e94ac65f7c85e9be90ffe24891454c210a6c447d7920c0abce94c8e7d04",
|
||||
".modelhub_state/market_intelligence.json": "82bfcb14b7ebecb3060ace53451841c8c504f0185fe3fdf3d0533ed05a27d86b",
|
||||
".modelhub_state/official_capabilities.json": "48b8616a328f11e04cfe380d9e7dff6df645408f428b702d67432aa8dfb0dc96",
|
||||
".modelhub_state/outcome_checkpoint.json": "38cf9c229e7987183287b05f5fd6388f23d331dd2c0ca11acc3faebe0eab7c89",
|
||||
".modelhub_state/queue_cleanup_latest.json": "f4afc5dbbf57b883780778e798518d72ad1770b7d52c2f9efa7ac7bca07813f7",
|
||||
".modelhub_state/recent_outcomes.jsonl": "192c324383a5677c649ff5959d4c94833cadd97e64612ac0cbb69a2d1c720104",
|
||||
".modelhub_state/recovery_active_tasks.jsonl": "63aaf1dab275a6a6144cec167deac14d5324368c393d71dd9d4c1b840dc52802",
|
||||
".modelhub_state/recovery_intents.jsonl": "e0bdf6d6facff5180b35f21bdcb0646b7d99a17aada8c864f2bcea09f655b51c",
|
||||
".modelhub_state/recovery_intents.jsonl": "dbd4eb203fa77c94ac47e7a1b458ba00d8073b9d87815989fdcddd7b9373bf57",
|
||||
".modelhub_state/routing_intelligence.json": "9cd6cab068960be9aa13a5d32b2c8fbec6d281eaae88659f6893211c5b0f0a7d",
|
||||
".modelhub_state/submission_exclusions.jsonl": "80315f370e9cc3cdadaf390e12573ef887794214779379c50b7b37b7631e8989",
|
||||
".modelhub_state/worker_crashes.jsonl": "07505b69e0b050da5f90f8b046a3069d3e1ca2574ca39896bc7d8a66293167f7",
|
||||
"ledger/submissions.jsonl": "91484b32049219522f891ed5b4a3e5d3d956a2a1c56261fc229f28c270f267be",
|
||||
"outcomes/submissions.jsonl": "b4dc5f70e132c83c53346ec03151b65c73e2fb82c7a712174bb0a780b8522f99"
|
||||
"outcomes/submissions.jsonl": "0151e4ee48024c1bb2b9540671c4d1976cc67bd4e093452dbf93ddba5a8b9c71"
|
||||
},
|
||||
"generation": 11334,
|
||||
"phase": "cycle",
|
||||
"generation": 11335,
|
||||
"phase": "intent",
|
||||
"schemaVersion": 1,
|
||||
"updatedAt": "2026-09-21T08:24:30.755531+00:00",
|
||||
"updatedAt": "2026-09-21T08:29:14.490172+00:00",
|
||||
"writerId": "8b35139af6674067a339a670222d4b67"
|
||||
}
|
||||
|
||||
@@ -484,8 +484,6 @@
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T11:30:54.600232+00:00", "modelId": "aisingapore/Qwen-SEA-LION-v4-32B-IT-8BIT", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34327989512, "estimatedRequiredGiB": 38.385, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 34346039951, "modelscopeLicense": "mit", "modelscopeParams": 32762123264, "modelscopeTags": ["license:mit", "model_type:qwen3", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 34346039951}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:24.235683+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969624", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T11:30:54.600185+00:00", "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 54864980000, "estimatedRequiredGiB": 61.361, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 54904930780, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 54904930780}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:24.241562+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969620", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:30:54.600294+00:00", "modelId": "XHToken/Spark-X2.5-4B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8224192408, "estimatedRequiredGiB": 9.209, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 8239718268, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4112079360, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 8239718268}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:24.199548+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969614", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:30:54.599869+00:00", "modelId": "XHToken/Spark-X2.5-1.7B", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3415340496, "estimatedRequiredGiB": 3.834, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 3430847031, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1707657216, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3430847031}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:31.991647+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969636", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:30:54.599934+00:00", "modelId": "XHToken/Spark-X2.5-4B-FP8", "modelProfile": {"architectures": ["Spark2_5ForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4450260320, "estimatedRequiredGiB": 4.991, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "spark2_5", "modelscopeFileSize": 4465562659, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:spark2_5", "library:safetensors", "library:", "task:text-generation", "custom_tag:llm", "custom_tag:sparkx2_5"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "fp8", "repositoryOnDiskBytes": 4465562659}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:31.935098+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969629", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T11:30:54.600352+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3.5-70B-R-NVFP4", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "1c0a4f69628b2a0ad37a6eee139a4c98afed7d8c93529a9ef42a90756e4e1940", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 42709318472, "estimatedRequiredGiB": 47.753, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 42728540075, "modelscopeLicense": "llama3.1", "modelscopeParams": 70553707056, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 42728540075}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:32.000203+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4969635", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:30:54.600422+00:00", "modelId": "Arain119/Sophia", "modelProfile": {"architectures": ["SophiaForCausalLM"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2226652224, "estimatedRequiredGiB": 7.28, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "sophia_hybrid", "modelscopeFileSize": 6513869659, "modelscopeLicense": "Apache License 2.0", "modelscopeParams": 1113293776, "modelscopeTags": ["license:Apache License 2.0", "model_type:sophia_hybrid", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6513869659}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:31.903313+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969626", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T11:30:54.600119+00:00", "modelId": "mlx-community/Devstral-Small-2-24B-Instruct-2512-OptiQ-4bit", "modelProfile": {"architectures": ["Mistral3ForConditionalGeneration"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 17367755740, "estimatedRequiredGiB": 19.429, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "mistral3", "modelscopeFileSize": 17385072379, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4929899520, "modelscopeTags": ["license:apache-2.0", "model_type:mistral3", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:mixed-precision", "custom_tag:4bit", "custom_tag:8bit", "custom_tag:optiq", "custom_tag:apple-silicon", "custom_tag:text-generation", "custom_tag:mistral", "custom_tag:mistral3", "custom_tag:devstral"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 17385072379}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T03:25:31.941721+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4969628", "taskType": "text-generation", "verifyResult": null}
|
||||
@@ -991,7 +989,7 @@
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T08:17:25.672503+00:00", "modelId": "neuralmagic/SmolLM-1.7B-Instruct-quantized.w8a16", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "d8df464812ff2fd3c46dc89b0fc6a96370e694fab8e52179f1ada2cdd0ea92b9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2014810120, "estimatedRequiredGiB": 2.256, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 2018195521, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1812039680, "modelscopeTags": ["license:apache-2.0", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2018195521}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T00:12:58.734996+00:00", "targetGpu": "Biren_166m", "taskId": "4990139", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T08:17:25.672517+00:00", "modelId": "neuralmagic/starcoder2-7b-FP8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7855165704, "estimatedRequiredGiB": 8.783, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 7858541078, "modelscopeLicense": "other", "modelscopeParams": 7400416256, "modelscopeTags": ["license:other", "model_type:starcoder2", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:fp8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 7858541078}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T00:12:58.738712+00:00", "targetGpu": "Biren_166m", "taskId": "4990140", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-21T08:22:03.151489+00:00", "modelId": "RedHatAI/starcoder2-3b-quantized.w8a8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4093158536, "estimatedRequiredGiB": 4.578, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 4096457676, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 3181366272, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4096457676}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T00:17:53.135774+00:00", "targetGpu": "Biren_166m", "taskId": "4990242", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "aisingapore/Llama-SEA-LION-v3-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.964, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16073634582, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16073634582}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T00:22:06.134784+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4990307", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-21T08:25:31.485626+00:00", "modelId": "aisingapore/Llama-SEA-LION-v3-8B", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16060556376, "estimatedRequiredGiB": 17.964, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 16073634582, "modelscopeLicense": "llama3.1", "modelscopeParams": 8030261248, "modelscopeTags": ["license:llama3.1", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16073634582}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-21T00:22:06.134784+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4990307", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "RedHatAI/gemma-2-2b-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6746735136, "estimatedRequiredGiB": 7.565, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 6768615358, "modelscopeLicense": "gemma", "modelscopeParams": 3204165888, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 6768615358}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T00:29:04.803696+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4990383", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "aisingapore/Gemma-SEA-LION-v4-27B-IT-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma3ForConditionalGeneration"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 29274355224, "estimatedRequiredGiB": 32.761, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma3", "modelscopeFileSize": 29314389092, "modelscopeLicense": "gemma", "modelscopeParams": 27432406640, "modelscopeTags": ["license:gemma", "model_type:gemma3", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 29314389092}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T00:29:04.738806+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4990382", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "neuralmagic/Llama-3.2-3B-Instruct-FP8", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2ebd1d3ee8ed952dd08e24120ad5d7342c9a018f55c478225a242b664721e8d0", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4395007696, "estimatedRequiredGiB": 4.922, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 4404161088, "modelscopeLicense": "llama3.2", "modelscopeParams": 3606752256, "modelscopeTags": ["license:llama3.2", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama", "custom_tag:llama-3", "custom_tag:neuralmagic", "custom_tag:llmcompressor"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4404161088}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-21T00:29:04.722455+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4990380", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user