state: generation 7161 (cycle)

This commit is contained in:
2026-09-14 11:52:06 +00:00
parent 59a9e23364
commit fab4d0bca3
9 changed files with 2107 additions and 2061 deletions

View File

@@ -1498,7 +1498,7 @@
"taskType": "text-generation" "taskType": "text-generation"
} }
}, },
"generatedAt": "2026-09-14T11:49:07.050592+00:00", "generatedAt": "2026-09-14T11:52:05.943793+00:00",
"summary": { "summary": {
"activeBlockCount": 75, "activeBlockCount": 75,
"byGpuFramework": { "byGpuFramework": {

View File

@@ -51,7 +51,7 @@
"4": { "4": {
"complete": false, "complete": false,
"lastError": "ModelHubAPIError: 系统错误", "lastError": "ModelHubAPIError: 系统错误",
"listingErrors": 520, "listingErrors": 521,
"nextPage": 1, "nextPage": 1,
"recordsScanned": 0, "recordsScanned": 0,
"uniqueRecords": 0 "uniqueRecords": 0
@@ -102,12 +102,12 @@
"cutoffAt": "2026-09-04T03:55:51.365685+00:00", "cutoffAt": "2026-09-04T03:55:51.365685+00:00",
"failureLogsInspected": 0, "failureLogsInspected": 0,
"mode": "incremental_decision_only", "mode": "incremental_decision_only",
"nextAccountIndex": 4, "nextAccountIndex": 5,
"recordsScanned": 0, "recordsScanned": 0,
"seenTaskIds": [], "seenTaskIds": [],
"startedAt": "2026-09-04T03:55:51.365685+00:00", "startedAt": "2026-09-04T03:55:51.365685+00:00",
"terminalRecords": 0, "terminalRecords": 0,
"uniqueRecords": 0, "uniqueRecords": 0,
"updatedAt": "2026-09-14T11:49:07.022929+00:00", "updatedAt": "2026-09-14T11:52:05.918421+00:00",
"version": 1 "version": 1
} }

View File

@@ -416,7 +416,7 @@
} }
}, },
"frameworkUpdatedAt": null, "frameworkUpdatedAt": null,
"generatedAt": "2026-09-14T11:47:20.052608+00:00", "generatedAt": "2026-09-14T11:50:21.383886+00:00",
"gpuStats": { "gpuStats": {
"Ascend_910-b3": { "Ascend_910-b3": {
"available": true, "available": true,

File diff suppressed because it is too large Load Diff

View File

@@ -1,6 +1,6 @@
{ {
"generatedAt": "2026-09-14T11:40:44.211337+00:00", "generatedAt": "2026-09-14T11:50:08.250891+00:00",
"lastSyncTime": "2026-09-14T11:40:43.317642+00:00", "lastSyncTime": "2026-09-14T11:50:08.015544+00:00",
"recentLimit": 300, "recentLimit": 300,
"report": { "report": {
"architectureCompatibilityBlocks": { "architectureCompatibilityBlocks": {
@@ -1932,28 +1932,28 @@
"unresolvedFailureCount": 32 "unresolvedFailureCount": 32
}, },
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation": { "Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation": {
"attributableFailureCount": 9, "attributableFailureCount": 10,
"decisionFailureRate": 0.8182, "decisionFailureRate": 0.8333,
"decisionSuccessRate": 0.1818, "decisionSuccessRate": 0.1667,
"decisionTotal": 11, "decisionTotal": 12,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 6, "ambiguous_runtime": 6,
"backend_operator": 3, "backend_operator": 3,
"framework_architecture_unsupported": 3, "framework_architecture_unsupported": 4,
"model_load": 1, "model_load": 1,
"runtime_memory": 2 "runtime_memory": 2
}, },
"failureCount": 15, "failureCount": 16,
"failureRate": 0.8824, "failureRate": 0.8889,
"framework": "vllm_0_17_0_corex_4_4_0", "framework": "vllm_0_17_0_corex_4_4_0",
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 2, "successCount": 2,
"successRate": 0.1176, "successRate": 0.1111,
"targetGpu": "Iluvatar_bi-150", "targetGpu": "Iluvatar_bi-150",
"taskType": "text-generation", "taskType": "text-generation",
"total": 17, "total": 18,
"unresolvedFailureCount": 6 "unresolvedFailureCount": 6
}, },
"Iluvatar_bi-150|vllm|text-generation": { "Iluvatar_bi-150|vllm|text-generation": {
@@ -2604,25 +2604,25 @@
"unresolvedFailureCount": 8 "unresolvedFailureCount": 8
}, },
"vllm_0_17_0_corex_4_4_0": { "vllm_0_17_0_corex_4_4_0": {
"attributableFailureCount": 9, "attributableFailureCount": 10,
"decisionFailureRate": 0.8182, "decisionFailureRate": 0.8333,
"decisionSuccessRate": 0.1818, "decisionSuccessRate": 0.1667,
"decisionTotal": 11, "decisionTotal": 12,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 6, "ambiguous_runtime": 6,
"backend_operator": 3, "backend_operator": 3,
"framework_architecture_unsupported": 3, "framework_architecture_unsupported": 4,
"model_load": 1, "model_load": 1,
"runtime_memory": 2 "runtime_memory": 2
}, },
"failureCount": 15, "failureCount": 16,
"failureRate": 0.8824, "failureRate": 0.8889,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 2, "successCount": 2,
"successRate": 0.1176, "successRate": 0.1111,
"total": 17, "total": 18,
"unresolvedFailureCount": 6 "unresolvedFailureCount": 6
}, },
"vllm_fix_tokenizer": { "vllm_fix_tokenizer": {
@@ -2668,7 +2668,7 @@
"unresolvedFailureCount": 1 "unresolvedFailureCount": 1
} }
}, },
"generatedAt": "2026-09-14T11:40:44.205598+00:00", "generatedAt": "2026-09-14T11:50:08.246900+00:00",
"gpuSummaries": { "gpuSummaries": {
"Ascend_910-b3": { "Ascend_910-b3": {
"attributableFailureCount": 33, "attributableFailureCount": 33,
@@ -2807,28 +2807,28 @@
"unresolvedFailureCount": 29 "unresolvedFailureCount": 29
}, },
"Iluvatar_bi-150": { "Iluvatar_bi-150": {
"attributableFailureCount": 36, "attributableFailureCount": 37,
"decisionFailureRate": 0.8182, "decisionFailureRate": 0.8222,
"decisionSuccessRate": 0.1818, "decisionSuccessRate": 0.1778,
"decisionTotal": 44, "decisionTotal": 45,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 14, "ambiguous_runtime": 14,
"backend_operator": 6, "backend_operator": 6,
"framework_architecture_unsupported": 17, "framework_architecture_unsupported": 18,
"model_load": 8, "model_load": 8,
"repository_structure": 1, "repository_structure": 1,
"runtime_memory": 4, "runtime_memory": 4,
"参数/模板问题": 2, "参数/模板问题": 2,
"验证失败": 30 "验证失败": 30
}, },
"failureCount": 82, "failureCount": 83,
"failureRate": 0.9111, "failureRate": 0.9121,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 8, "successCount": 8,
"successRate": 0.0889, "successRate": 0.0879,
"total": 90, "total": 91,
"unresolvedFailureCount": 46 "unresolvedFailureCount": 46
}, },
"Iluvatar_mrv-100": { "Iluvatar_mrv-100": {
@@ -3547,6 +3547,29 @@
"total": 4, "total": 4,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|nanbeige|compressed-tensors": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm_0_17_0_corex_4_4_0",
"modelType": "nanbeige",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Iluvatar_bi-150",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|phi3|compressed-tensors": { "Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|phi3|compressed-tensors": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"decisionFailureRate": 0.0, "decisionFailureRate": 0.0,
@@ -6254,6 +6277,30 @@
"total": 2, "total": 2,
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|nanbeige|compressed-tensors|32": {
"attributableFailureCount": 1,
"decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0,
"decisionTotal": 1,
"failureBreakdown": {
"framework_architecture_unsupported": 1
},
"failureCount": 1,
"failureRate": 1.0,
"framework": "vllm_0_17_0_corex_4_4_0",
"loadSizeLog2Bucket": 32,
"modelType": "nanbeige",
"pendingCount": 0,
"pendingRate": 0.0,
"platformFailureCount": 0,
"quantizationMethod": "compressed-tensors",
"successCount": 0,
"successRate": 0.0,
"targetGpu": "Iluvatar_bi-150",
"taskType": "text-generation",
"total": 1,
"unresolvedFailureCount": 0
},
"Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|phi3|compressed-tensors|33": { "Iluvatar_bi-150|vllm_0_17_0_corex_4_4_0|text-generation|phi3|compressed-tensors|33": {
"attributableFailureCount": 0, "attributableFailureCount": 0,
"decisionFailureRate": 0.0, "decisionFailureRate": 0.0,
@@ -8195,17 +8242,17 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
} }
}, },
"terminalRecords": 1828, "terminalRecords": 1829,
"totalRecords": 1921, "totalRecords": 1922,
"totals": { "totals": {
"attributableFailureCount": 548, "attributableFailureCount": 549,
"decisionFailureRate": 0.9013, "decisionFailureRate": 0.9015,
"decisionSuccessRate": 0.0987, "decisionSuccessRate": 0.0985,
"decisionTotal": 608, "decisionTotal": 609,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 432, "ambiguous_runtime": 432,
"backend_operator": 43, "backend_operator": 43,
"framework_architecture_unsupported": 348, "framework_architecture_unsupported": 349,
"memory_capacity": 12, "memory_capacity": 12,
"model_load": 76, "model_load": 76,
"platform_infrastructure": 6, "platform_infrastructure": 6,
@@ -8215,14 +8262,14 @@
"参数/模板问题": 109, "参数/模板问题": 109,
"验证失败": 673 "验证失败": 673
}, },
"failureCount": 1768, "failureCount": 1769,
"failureRate": 0.9672, "failureRate": 0.9672,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 6, "platformFailureCount": 6,
"successCount": 60, "successCount": 60,
"successRate": 0.0328, "successRate": 0.0328,
"total": 1828, "total": 1829,
"unresolvedFailureCount": 1214 "unresolvedFailureCount": 1214
}, },
"warnings": [ "warnings": [
@@ -8261,6 +8308,6 @@
] ]
}, },
"storageMode": "decision_state_only", "storageMode": "decision_state_only",
"summarizedRecords": 1921, "summarizedRecords": 1922,
"version": 1 "version": 1
} }

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -1,24 +1,24 @@
{ {
"agentVersion": "2026.09.04.4", "agentVersion": "2026.09.04.4",
"checksums": { "checksums": {
".modelhub_state/architecture_compatibility_blacklist.json": "3c0f425ac283493a84dbd36eaafd9ce8d8ff8ddb01a7d696ef9e2390d29339d8", ".modelhub_state/architecture_compatibility_blacklist.json": "5ae6a3a812c235a133e4033a73c37450430369dc434ab0559c5c92615a22fa9d",
".modelhub_state/architecture_history_backfill.json": "49758f5ccc466e552246a2ad7040bceb591b3e83e3b14b6b97a896d3cae91e35", ".modelhub_state/architecture_history_backfill.json": "744cd33138f8ad832af925d194cc4eca26bd25aa275646fd860485174e5f8f77",
".modelhub_state/market_intelligence.json": "a2fd6d04d93c0a2bff2e46076172eb84117cfcb102b3f85ba6cf0ef9e6d1e7b6", ".modelhub_state/market_intelligence.json": "7ffc9b7fc6a034d625d3d0c8000b8c24fc37a1314c3b72e0e87f32cf550047bd",
".modelhub_state/official_capabilities.json": "578a488738de6f5c417df9b3e67a412311864b37c4bd51b91dcb422ad43fe6a5", ".modelhub_state/official_capabilities.json": "bb8be5ae8eb012489df557107a5215b01ee8b1c6258c6d6a44bef94ba679e0e5",
".modelhub_state/outcome_checkpoint.json": "c69203b6c4bb7a21971ea908c98542e5307a01bf2993e04d1a26d942bf379cc3", ".modelhub_state/outcome_checkpoint.json": "2e82844a2a2bd49b6d0a116552cc8a127cb1c4874c5a69630c083650f6974d40",
".modelhub_state/queue_cleanup_latest.json": "e5d9b334b2784235be02767778d1d31aa5d889dbe0cdb9bf65bb77a097c78531", ".modelhub_state/queue_cleanup_latest.json": "e5d9b334b2784235be02767778d1d31aa5d889dbe0cdb9bf65bb77a097c78531",
".modelhub_state/recent_outcomes.jsonl": "fec83e03ae4b3a0c872fb87d8c96b9dfa135fd131a6af08858404e71a6b4c4c9", ".modelhub_state/recent_outcomes.jsonl": "fec83e03ae4b3a0c872fb87d8c96b9dfa135fd131a6af08858404e71a6b4c4c9",
".modelhub_state/recovery_active_tasks.jsonl": "c2d755696157fa62cc28d9ef7fbc855d80b2661e24c93f3a459c8c68bda95071", ".modelhub_state/recovery_active_tasks.jsonl": "4a8e55f87c626a46336fdeace0d069d3e9703dc9aab06d55c2fd8374003f7f37",
".modelhub_state/recovery_intents.jsonl": "6f4afcb3eb70766911990c0563a7a1f10667b60f02bfb7e64e4c2cb8a55cd997", ".modelhub_state/recovery_intents.jsonl": "10c05804bad9fb6751ae3c3d3568e959bfb73ca96803e30c45111771f8160a03",
".modelhub_state/routing_intelligence.json": "350794262038f5417f2c2055e35c53a6da870601a90b6c6b470484a8808e8764", ".modelhub_state/routing_intelligence.json": "350794262038f5417f2c2055e35c53a6da870601a90b6c6b470484a8808e8764",
".modelhub_state/submission_exclusions.jsonl": "632739c6cd6a2dd2237572ba18e6a390c5aa37ce56a18a2c1341d78569421552", ".modelhub_state/submission_exclusions.jsonl": "632739c6cd6a2dd2237572ba18e6a390c5aa37ce56a18a2c1341d78569421552",
".modelhub_state/worker_crashes.jsonl": "53505d335313446ff7094bcb58658b561de55799302fa7bfb49e418e414eebad", ".modelhub_state/worker_crashes.jsonl": "53505d335313446ff7094bcb58658b561de55799302fa7bfb49e418e414eebad",
"ledger/submissions.jsonl": "534da75337c614c3a3b5ecc6a07cf615533851ea1ac60f312bfbcce339099da3", "ledger/submissions.jsonl": "534da75337c614c3a3b5ecc6a07cf615533851ea1ac60f312bfbcce339099da3",
"outcomes/submissions.jsonl": "0ea5146c07efcd0902da1586dfeb46b7927efdf4b73e3cc463b6531f8378342e" "outcomes/submissions.jsonl": "f870c1f40446ecdb1dc7ada1d9077e285b37ce74b903a7b46bd8e5d51e75cef4"
}, },
"generation": 7160, "generation": 7161,
"phase": "cycle", "phase": "cycle",
"schemaVersion": 1, "schemaVersion": 1,
"updatedAt": "2026-09-14T11:49:07.155302+00:00", "updatedAt": "2026-09-14T11:52:06.644099+00:00",
"writerId": "69701e73196b4ac2b3fad0f412335168" "writerId": "69701e73196b4ac2b3fad0f412335168"
} }

View File

@@ -117,7 +117,6 @@
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-05T22:00:13.693287+00:00", "modelId": "FlagRelease/DeepSeek-R1-Distill-Qwen-1.5B-mthreads-FlagOS", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3554214621, "estimatedRequiredGiB": 3.984, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 3564406686, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1777088000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3564406686}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T13:55:22.931577+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4650076", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-05T22:00:13.693287+00:00", "modelId": "FlagRelease/DeepSeek-R1-Distill-Qwen-1.5B-mthreads-FlagOS", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3554214621, "estimatedRequiredGiB": 3.984, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 3564406686, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1777088000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 3564406686}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T13:55:22.931577+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4650076", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T22:35:47.104165+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "99fcfb5673a881c74d82451063f2e7a90a328c33bce1e63a9d0bf3ccffd99170", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T14:30:50.543106+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4650543", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T22:35:47.104165+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "99fcfb5673a881c74d82451063f2e7a90a328c33bce1e63a9d0bf3ccffd99170", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T14:30:50.543106+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4650543", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T22:44:15.950234+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "05c5a771e8ede0f61c61c1696314e8cb295d87d07704eedc68e43f7483931e45", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T14:39:36.525018+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4650654", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-09-05T22:44:15.950234+00:00", "modelId": "LiquidAI/LFM2.5-8B-A1B-DSpark-GGUF", "modelProfile": {"architectures": [], "configFingerprint": "05c5a771e8ede0f61c61c1696314e8cb295d87d07704eedc68e43f7483931e45", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 356491104, "estimatedRequiredGiB": 1.363, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 1219879026, "modelscopeLicense": "other", "modelscopeParams": 327707521, "modelscopeTags": ["license:other", "library:gguf", "library:pytorch", "task:text-generation", "custom_tag:speculative-decoding", "custom_tag:dspark", "custom_tag:lfm2", "custom_tag:draft-model", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 1219879026}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T14:39:36.525018+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4650654", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-05T23:08:33.934658+00:00", "modelId": "nanbeige/Nanbeige4.2-3B-GPTQ-Int8", "modelProfile": {"architectures": ["NanbeigeForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5192364968, "estimatedRequiredGiB": 5.827, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "nanbeige", "modelscopeFileSize": 5214325469, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4169800704, "modelscopeTags": ["license:apache-2.0", "model_type:nanbeige", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:llm", "custom_tag:nanbeige"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 5214325469}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T15:03:56.182791+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4650921", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113639+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-DPO-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737317601, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737317601}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.385517+00:00", "targetGpu": "Vastai_va16", "taskId": "4651796", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113639+00:00", "modelId": "solidrust/Llama-3-8B-Instruct-DPO-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5727954960, "estimatedRequiredGiB": 6.412, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 5737317601, "modelscopeLicense": "other", "modelscopeParams": 8030269440, "modelscopeTags": ["license:other", "model_type:llama", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:axolotl", "custom_tag:finetune", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 5737317601}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.385517+00:00", "targetGpu": "Vastai_va16", "taskId": "4651796", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113688+00:00", "modelId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15338408461, "estimatedRequiredGiB": 17.168, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 15362064483, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7669052656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:exl3", "custom_tag:exllamav3", "custom_tag:quantization", "custom_tag:long-context", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "exl3", "repositoryOnDiskBytes": 15362064483}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.391817+00:00", "targetGpu": "Vastai_va16", "taskId": "4651794", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113688+00:00", "modelId": "Mia-AiLab/Qwen3.8-27B-EXL3-3.5bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15338408461, "estimatedRequiredGiB": 17.168, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 15362064483, "modelscopeLicense": "apache-2.0", "modelscopeParams": 7669052656, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:exl3", "custom_tag:exllamav3", "custom_tag:quantization", "custom_tag:long-context", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "exl3", "repositoryOnDiskBytes": 15362064483}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.391817+00:00", "targetGpu": "Vastai_va16", "taskId": "4651794", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113630+00:00", "modelId": "solidrust/Llama-3-11B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7541230432, "estimatedRequiredGiB": 8.438, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 7550622334, "modelscopeLicense": "other", "modelscopeParams": 11520053248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 7550622334}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.389388+00:00", "targetGpu": "Vastai_va16", "taskId": "4651792", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-06T00:23:33.113630+00:00", "modelId": "solidrust/Llama-3-11B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "cbe5569cdb6c42553436b837dd5b8f0cfa03334642a56f4ff0137f952f8eb84b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7541230432, "estimatedRequiredGiB": 8.438, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 7550622334, "modelscopeLicense": "other", "modelscopeParams": 11520053248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 7550622334}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-05T16:20:38.389388+00:00", "targetGpu": "Vastai_va16", "taskId": "4651792", "taskType": "text-generation", "verifyResult": null}