state: generation 8379 (cycle)

This commit is contained in:
2026-09-16 20:33:16 +00:00
parent bd7cdb0d20
commit f40cacf5a9
9 changed files with 2028 additions and 2025 deletions

View File

@@ -1714,7 +1714,7 @@
"taskType": "text-generation" "taskType": "text-generation"
} }
}, },
"generatedAt": "2026-09-16T20:29:46.954326+00:00", "generatedAt": "2026-09-16T20:33:16.119690+00:00",
"summary": { "summary": {
"activeBlockCount": 83, "activeBlockCount": 83,
"byGpuFramework": { "byGpuFramework": {

View File

@@ -51,7 +51,7 @@
"4": { "4": {
"complete": false, "complete": false,
"lastError": "ModelHubAPIError: 系统错误", "lastError": "ModelHubAPIError: 系统错误",
"listingErrors": 611, "listingErrors": 612,
"nextPage": 1, "nextPage": 1,
"recordsScanned": 0, "recordsScanned": 0,
"uniqueRecords": 0 "uniqueRecords": 0
@@ -102,12 +102,12 @@
"cutoffAt": "2026-09-04T03:55:51.365685+00:00", "cutoffAt": "2026-09-04T03:55:51.365685+00:00",
"failureLogsInspected": 0, "failureLogsInspected": 0,
"mode": "incremental_decision_only", "mode": "incremental_decision_only",
"nextAccountIndex": 4, "nextAccountIndex": 5,
"recordsScanned": 0, "recordsScanned": 0,
"seenTaskIds": [], "seenTaskIds": [],
"startedAt": "2026-09-04T03:55:51.365685+00:00", "startedAt": "2026-09-04T03:55:51.365685+00:00",
"terminalRecords": 0, "terminalRecords": 0,
"uniqueRecords": 0, "uniqueRecords": 0,
"updatedAt": "2026-09-16T20:29:46.930358+00:00", "updatedAt": "2026-09-16T20:33:16.094696+00:00",
"version": 1 "version": 1
} }

View File

@@ -1,8 +1,8 @@
{ {
"communityAttemptedAt": "2026-09-16T20:13:16.094732+00:00", "communityAttemptedAt": "2026-09-16T20:30:57.508604+00:00",
"communityError": null, "communityError": null,
"communitySample": {}, "communitySample": {},
"communityUpdatedAt": "2026-09-16T20:13:16.094732+00:00", "communityUpdatedAt": "2026-09-16T20:30:57.508604+00:00",
"frameworkAttemptedAt": "2026-09-16T14:37:17.574684+00:00", "frameworkAttemptedAt": "2026-09-16T14:37:17.574684+00:00",
"frameworkError": null, "frameworkError": null,
"frameworkStats": { "frameworkStats": {
@@ -416,7 +416,7 @@
} }
}, },
"frameworkUpdatedAt": null, "frameworkUpdatedAt": null,
"generatedAt": "2026-09-16T20:28:13.557549+00:00", "generatedAt": "2026-09-16T20:30:57.508604+00:00",
"gpuStats": { "gpuStats": {
"Ascend_910-b3": { "Ascend_910-b3": {
"available": true, "available": true,

File diff suppressed because it is too large Load Diff

View File

@@ -1,6 +1,6 @@
{ {
"generatedAt": "2026-09-16T20:22:35.065787+00:00", "generatedAt": "2026-09-16T20:30:49.878292+00:00",
"lastSyncTime": "2026-09-16T20:22:35.024646+00:00", "lastSyncTime": "2026-09-16T20:30:47.710570+00:00",
"recentLimit": 300, "recentLimit": 300,
"report": { "report": {
"architectureCompatibilityBlocks": { "architectureCompatibilityBlocks": {
@@ -2692,16 +2692,17 @@
"unresolvedFailureCount": 29 "unresolvedFailureCount": 29
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation": {
"attributableFailureCount": 6, "attributableFailureCount": 7,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 6, "decisionTotal": 7,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 3, "ambiguous_runtime": 3,
"framework_architecture_unsupported": 6, "framework_architecture_unsupported": 6,
"runtime_memory": 1,
"参数/模板问题": 8 "参数/模板问题": 8
}, },
"failureCount": 17, "failureCount": 18,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm-patch-tokenizer", "framework": "vllm-patch-tokenizer",
"pendingCount": 0, "pendingCount": 0,
@@ -2711,7 +2712,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "hygon_k100-ai", "targetGpu": "hygon_k100-ai",
"taskType": "text-generation", "taskType": "text-generation",
"total": 17, "total": 18,
"unresolvedFailureCount": 11 "unresolvedFailureCount": 11
}, },
"hygon_k100-ai|vllm|text-generation": { "hygon_k100-ai|vllm|text-generation": {
@@ -2855,23 +2856,24 @@
"unresolvedFailureCount": 34 "unresolvedFailureCount": 34
}, },
"vllm-patch-tokenizer": { "vllm-patch-tokenizer": {
"attributableFailureCount": 6, "attributableFailureCount": 7,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 6, "decisionTotal": 7,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 3, "ambiguous_runtime": 3,
"framework_architecture_unsupported": 6, "framework_architecture_unsupported": 6,
"runtime_memory": 1,
"参数/模板问题": 8 "参数/模板问题": 8
}, },
"failureCount": 17, "failureCount": 18,
"failureRate": 1.0, "failureRate": 1.0,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 0, "successCount": 0,
"successRate": 0.0, "successRate": 0.0,
"total": 17, "total": 18,
"unresolvedFailureCount": 11 "unresolvedFailureCount": 11
}, },
"vllm_0_17_0_corex_4_4_0": { "vllm_0_17_0_corex_4_4_0": {
@@ -2941,7 +2943,7 @@
"unresolvedFailureCount": 5 "unresolvedFailureCount": 5
} }
}, },
"generatedAt": "2026-09-16T20:22:35.059022+00:00", "generatedAt": "2026-09-16T20:30:49.872495+00:00",
"gpuSummaries": { "gpuSummaries": {
"Ascend_910-b3": { "Ascend_910-b3": {
"attributableFailureCount": 49, "attributableFailureCount": 49,
@@ -3254,28 +3256,28 @@
"unresolvedFailureCount": 213 "unresolvedFailureCount": 213
}, },
"hygon_k100-ai": { "hygon_k100-ai": {
"attributableFailureCount": 60, "attributableFailureCount": 61,
"decisionFailureRate": 1.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 60, "decisionTotal": 61,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 27, "ambiguous_runtime": 27,
"framework_architecture_unsupported": 48, "framework_architecture_unsupported": 48,
"memory_capacity": 1, "memory_capacity": 1,
"model_load": 5, "model_load": 5,
"repository_structure": 4, "repository_structure": 4,
"runtime_memory": 2, "runtime_memory": 3,
"参数/模板问题": 10, "参数/模板问题": 10,
"验证失败": 27 "验证失败": 27
}, },
"failureCount": 124, "failureCount": 125,
"failureRate": 1.0, "failureRate": 1.0,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 0, "platformFailureCount": 0,
"successCount": 0, "successCount": 0,
"successRate": 0.0, "successRate": 0.0,
"total": 124, "total": 125,
"unresolvedFailureCount": 64 "unresolvedFailureCount": 64
} }
}, },
@@ -6095,14 +6097,15 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|llama|awq": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation|llama|awq": {
"attributableFailureCount": 0, "attributableFailureCount": 1,
"decisionFailureRate": 0.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 0, "decisionTotal": 1,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 3 "ambiguous_runtime": 3,
"runtime_memory": 1
}, },
"failureCount": 3, "failureCount": 4,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm-patch-tokenizer", "framework": "vllm-patch-tokenizer",
"modelType": "llama", "modelType": "llama",
@@ -6114,7 +6117,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "hygon_k100-ai", "targetGpu": "hygon_k100-ai",
"taskType": "text-generation", "taskType": "text-generation",
"total": 3, "total": 4,
"unresolvedFailureCount": 3 "unresolvedFailureCount": 3
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|nanbeige|compressed-tensors": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation|nanbeige|compressed-tensors": {
@@ -11086,14 +11089,15 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|llama|awq|32": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation|llama|awq|32": {
"attributableFailureCount": 0, "attributableFailureCount": 1,
"decisionFailureRate": 0.0, "decisionFailureRate": 1.0,
"decisionSuccessRate": 0.0, "decisionSuccessRate": 0.0,
"decisionTotal": 0, "decisionTotal": 1,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 3 "ambiguous_runtime": 3,
"runtime_memory": 1
}, },
"failureCount": 3, "failureCount": 4,
"failureRate": 1.0, "failureRate": 1.0,
"framework": "vllm-patch-tokenizer", "framework": "vllm-patch-tokenizer",
"loadSizeLog2Bucket": 32, "loadSizeLog2Bucket": 32,
@@ -11106,7 +11110,7 @@
"successRate": 0.0, "successRate": 0.0,
"targetGpu": "hygon_k100-ai", "targetGpu": "hygon_k100-ai",
"taskType": "text-generation", "taskType": "text-generation",
"total": 3, "total": 4,
"unresolvedFailureCount": 3 "unresolvedFailureCount": 3
}, },
"hygon_k100-ai|vllm-patch-tokenizer|text-generation|nanbeige|compressed-tensors|32": { "hygon_k100-ai|vllm-patch-tokenizer|text-generation|nanbeige|compressed-tensors|32": {
@@ -11230,13 +11234,13 @@
"unresolvedFailureCount": 0 "unresolvedFailureCount": 0
} }
}, },
"terminalRecords": 2046, "terminalRecords": 2047,
"totalRecords": 2149, "totalRecords": 2150,
"totals": { "totals": {
"attributableFailureCount": 663, "attributableFailureCount": 664,
"decisionFailureRate": 0.8911, "decisionFailureRate": 0.8913,
"decisionSuccessRate": 0.1089, "decisionSuccessRate": 0.1087,
"decisionTotal": 744, "decisionTotal": 745,
"failureBreakdown": { "failureBreakdown": {
"ambiguous_runtime": 495, "ambiguous_runtime": 495,
"backend_operator": 45, "backend_operator": 45,
@@ -11245,19 +11249,19 @@
"model_load": 85, "model_load": 85,
"platform_infrastructure": 10, "platform_infrastructure": 10,
"repository_structure": 8, "repository_structure": 8,
"runtime_memory": 17, "runtime_memory": 18,
"tokenizer_compatibility": 60, "tokenizer_compatibility": 60,
"参数/模板问题": 124, "参数/模板问题": 124,
"验证失败": 673 "验证失败": 673
}, },
"failureCount": 1965, "failureCount": 1966,
"failureRate": 0.9604, "failureRate": 0.9604,
"pendingCount": 0, "pendingCount": 0,
"pendingRate": 0.0, "pendingRate": 0.0,
"platformFailureCount": 10, "platformFailureCount": 10,
"successCount": 81, "successCount": 81,
"successRate": 0.0396, "successRate": 0.0396,
"total": 2046, "total": 2047,
"unresolvedFailureCount": 1292 "unresolvedFailureCount": 1292
}, },
"warnings": [ "warnings": [
@@ -11268,11 +11272,11 @@
"GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。", "GPU Biren_166m 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。", "GPU Iluvatar_bi-150 本地统计失败率偏高≥50%),建议重点关注。",
"GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。", "GPU hygon_k100-ai 本地统计失败率偏高≥50%),建议重点关注。",
"GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。", "GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。", "GPU Cambricon_mlu-370-x4 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。", "GPU Ascend_910-b3 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。", "GPU Cambricon_mlu-370-x8 本地统计失败率偏高≥50%),建议重点关注。",
"GPU Vastai_va16 本地统计失败率偏高≥50%),建议重点关注。", "GPU MetaX_c-500 本地统计失败率偏高≥50%),建议重点关注。",
"组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Biren_166m|unknown|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Biren_166m|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Biren_166m|vllm|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
"组合 Vastai_va16|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。", "组合 Vastai_va16|vllm_fix_tokenizer|text-generation 近期失败集中,建议降低该 GPU+框架的提交优先级。",
@@ -11297,6 +11301,6 @@
] ]
}, },
"storageMode": "decision_state_only", "storageMode": "decision_state_only",
"summarizedRecords": 2149, "summarizedRecords": 2150,
"version": 1 "version": 1
} }

File diff suppressed because it is too large Load Diff

File diff suppressed because it is too large Load Diff

View File

@@ -1,24 +1,24 @@
{ {
"agentVersion": "2026.09.04.4", "agentVersion": "2026.09.04.4",
"checksums": { "checksums": {
".modelhub_state/architecture_compatibility_blacklist.json": "3d69337fc6f2d99fd5f6e988ecd2806fb1d5fcb4cceee8d14e45761b5861bb24", ".modelhub_state/architecture_compatibility_blacklist.json": "99f7db8646f3b0d7356f464231d0093a7ce8ad11189c7677e599a25dbb05cc88",
".modelhub_state/architecture_history_backfill.json": "7038ed8604e3323053651a805f901688ce59211196474fb64a81b7c2f9048356", ".modelhub_state/architecture_history_backfill.json": "b9102bab7bb0159808a3e90bbdc4074a446b841b65d1533f747fdcf5ad506981",
".modelhub_state/market_intelligence.json": "d23fcc7ef7ab21d0ed2f1636afb8ae0d13b77e716fbad7efe6d4cd4748f7e0b9", ".modelhub_state/market_intelligence.json": "320bd4eee373b6cd4afd85b991852ef4a742be0ea7dc479ce6c1c8a5154f2564",
".modelhub_state/official_capabilities.json": "e4a6dbacb591825ea47fb2037c41150e900687eb65b120a8de8cb2f073905e3c", ".modelhub_state/official_capabilities.json": "d2bd3102d9d2ec37adf508351e57034701987c1e7e28dedb1a69fd4450e96cd1",
".modelhub_state/outcome_checkpoint.json": "441d219971455cd268ef05da285daac632c47b6abfbc86a66b83f44d806fb400", ".modelhub_state/outcome_checkpoint.json": "5bfe9462aea424dcabeab62683627d30fce358c88e8c9b192398fad5fb86eaac",
".modelhub_state/queue_cleanup_latest.json": "4b07ec677c9f7a1cef8978746f6e057a9979372430d32644646c3bd44e0a0546", ".modelhub_state/queue_cleanup_latest.json": "4b07ec677c9f7a1cef8978746f6e057a9979372430d32644646c3bd44e0a0546",
".modelhub_state/recent_outcomes.jsonl": "0956a3bdee9f38177b96f0df07e4547a84aaa98a6fcec842c84ec4d32f68befc", ".modelhub_state/recent_outcomes.jsonl": "0956a3bdee9f38177b96f0df07e4547a84aaa98a6fcec842c84ec4d32f68befc",
".modelhub_state/recovery_active_tasks.jsonl": "ce8de5b8e317aa924d465209a9c812820b8afdabfe4f53ee837dba3de1a63cc7", ".modelhub_state/recovery_active_tasks.jsonl": "0eab2f753bb911b9c29a44acc5b2053e2526d3246eeee7419bd0f49069c2f242",
".modelhub_state/recovery_intents.jsonl": "5284ad4aa360ef4a07c503451f3ec684172a94a7333c57dd710a867a50347917", ".modelhub_state/recovery_intents.jsonl": "bc96c1f1799aac48dad56f79ea959d2b9669707c77a7da9c15bd18224d46ca8b",
".modelhub_state/routing_intelligence.json": "008e6b54a9b7e1bc15577d1efc269953cea31d9de94d820a59dc3a61a5e8019c", ".modelhub_state/routing_intelligence.json": "008e6b54a9b7e1bc15577d1efc269953cea31d9de94d820a59dc3a61a5e8019c",
".modelhub_state/submission_exclusions.jsonl": "632739c6cd6a2dd2237572ba18e6a390c5aa37ce56a18a2c1341d78569421552", ".modelhub_state/submission_exclusions.jsonl": "632739c6cd6a2dd2237572ba18e6a390c5aa37ce56a18a2c1341d78569421552",
".modelhub_state/worker_crashes.jsonl": "b41d1dda7bab11bedd96de40fecd2d416e47396901b3283f48a56de5d6348983", ".modelhub_state/worker_crashes.jsonl": "b41d1dda7bab11bedd96de40fecd2d416e47396901b3283f48a56de5d6348983",
"ledger/submissions.jsonl": "f2e9c2422a6cd9a86afc55b3350a130222ff01ee5666806f2ac0b931cbf00c81", "ledger/submissions.jsonl": "f2e9c2422a6cd9a86afc55b3350a130222ff01ee5666806f2ac0b931cbf00c81",
"outcomes/submissions.jsonl": "595b6909501b46c75a0cef7ef56fbe7d4b626b0424d118b4e21267b21a073512" "outcomes/submissions.jsonl": "06b9af1dc6273ad060a58cc1d5726e0fddd22a6ca40f6517a3ef05325e857eb3"
}, },
"generation": 8378, "generation": 8379,
"phase": "cycle", "phase": "cycle",
"schemaVersion": 1, "schemaVersion": 1,
"updatedAt": "2026-09-16T20:29:47.056050+00:00", "updatedAt": "2026-09-16T20:33:16.617510+00:00",
"writerId": "eccb3e0018f640d19e578c271a207b5c" "writerId": "eccb3e0018f640d19e578c271a207b5c"
} }

View File

@@ -16,7 +16,6 @@
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-04T23:24:55.604425+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "manufacturer_spec", "sourceUrl": "https://www.birentech.com/news/id6rz98v3obczy77cmxzfgk3/"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 22125413673, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T15:17:04.566515+00:00", "targetGpu": "Biren_166m", "taskId": "4629091", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-04T23:24:55.604425+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "d341b8002cf39d3299fd429479e8c8e954583e49d4fec5a640ddc0f8f0cbd054", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "manufacturer_spec", "sourceUrl": "https://www.birentech.com/news/id6rz98v3obczy77cmxzfgk3/"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 22125413673, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T15:17:04.566515+00:00", "targetGpu": "Biren_166m", "taskId": "4629091", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-04T23:34:40.935412+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "published_card_spec", "sourceUrl": "https://aclanthology.org/2025.emnlp-main.1630.pdf"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T15:32:11.647812+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4629294", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-04T23:34:40.935412+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "published_card_spec", "sourceUrl": "https://aclanthology.org/2025.emnlp-main.1630.pdf"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T15:32:11.647812+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4629294", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-04T23:34:40.935439+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:15"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 22125413673, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T15:32:11.649244+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4629295", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-04T23:34:40.935439+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:15"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 22125413673, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T15:32:11.649244+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4629295", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-04T23:48:19.100125+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:15"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T15:47:16.843035+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4629532", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T00:08:56.210086+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://docs.mthreads.com/s4000/s4000-doc-online/product_specifications/"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T16:05:38.973026+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4629768", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T00:08:56.210086+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "manufacturer_spec", "sourceUrl": "https://docs.mthreads.com/s4000/s4000-doc-online/product_specifications/"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T16:05:38.973026+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4629768", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T00:12:22.613750+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:29"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 44221632439, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T16:10:19.951781+00:00", "targetGpu": "MetaX_c-500", "taskId": "4629851", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-05T00:12:22.613750+00:00", "modelId": "amd/gpt-oss-20b-BF16-w8a8-llmcompressor", "modelProfile": {"architectures": ["GptOssForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 44193251544, "estimatedRequiredGiB": 49.422, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:29"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gpt_oss", "modelscopeFileSize": 44221632439, "modelscopeLicense": "apache-2.0", "modelscopeParams": 20914757184, "modelscopeTags": ["license:apache-2.0", "model_type:gpt_oss", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:quantized", "custom_tag:int8", "custom_tag:w8a8", "custom_tag:dynamic-quantization", "custom_tag:8-bit", "custom_tag:llm-compressor", "custom_tag:compressed-tensors", "custom_tag:zendnn", "custom_tag:amd", "custom_tag:cpu-inference"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 44221632439}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T16:10:19.951781+00:00", "targetGpu": "MetaX_c-500", "taskId": "4629851", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T00:55:08.690909+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T16:50:16.092995+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4630839", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-05T00:55:08.690909+00:00", "modelId": "solidrust/Llama-3-13B-Instruct-v0.1-AWQ", "modelProfile": {"architectures": ["LlamaForCausalLM"], "configFingerprint": "2e9cebc01b20081f0839fa18fe9ca4f7d059509f235699027201fe328177b304", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8447876488, "estimatedRequiredGiB": 9.452, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "modelhub_preflight_oom:36"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "llama", "modelscopeFileSize": 8457284244, "modelscopeLicense": "other", "modelscopeParams": 13264949248, "modelscopeTags": ["license:other", "model_type:llama", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:4-bit", "custom_tag:AWQ", "custom_tag:text-generation", "custom_tag:autotrain_compatible", "custom_tag:endpoints_compatible", "custom_tag:mergekit", "custom_tag:merge", "custom_tag:facebook", "custom_tag:meta", "custom_tag:pytorch", "custom_tag:llama", "custom_tag:llama-3"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "awq", "repositoryOnDiskBytes": 8457284244}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-04T16:50:16.092995+00:00", "targetGpu": "Sunrise_pt-200-x1", "taskId": "4630839", "taskType": "text-generation", "verifyResult": null}