state: generation 8329 (intent)

This commit is contained in:
2026-09-16 18:09:58 +00:00
parent 19800640e8
commit 6ac2e1155f
8 changed files with 832 additions and 832 deletions

View File

@@ -623,7 +623,7 @@
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-16T17:25:59.446149+00:00", "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "533968aac3b9c90272dfbe2866791f782e077f0dceee5b6ffffce48ba6674a10", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-16T09:20:28.174669+00:00", "targetGpu": "MetaX_c-500", "taskId": "4889839", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-mlu", "lastSyncTime": "2026-09-16T17:41:50.112754+00:00", "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "36ec94e289a8da5fd70b98a2f401fc5159740daf9e2a076da4fd490f34bcbc0b", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 24.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 24.0, "maximumRepositorySizeGiB": 20.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-16T09:37:59.405066+00:00", "targetGpu": "Cambricon_mlu-370-x4", "taskId": "4890080", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": "2026-09-16T17:50:41.534321+00:00", "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-16T09:49:03.662873+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4890217", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T09:59:53.000301+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4890315", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-16T18:07:34.714203+00:00", "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "71d6e16091ba0289f01acffc8e39bb8c25703488dc49b9e390e08951fe1f2b33", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-16T09:59:53.000301+00:00", "targetGpu": "Iluvatar_mrv-100", "taskId": "4890315", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "a494b1667fab53609b94cc853db1a12c8d3b9fcf98795ecbab3dd5eb3a3c77f6", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T10:17:32.995862+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4890520", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": null, "modelId": "Edge0/Edge0-8B-A1B-preview", "modelProfile": {"architectures": ["BailingMoeV3ForCausalLM"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4563913288, "estimatedRequiredGiB": 5.138, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": null, "modelscopeFileSize": 4597698253, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1268239776, "modelscopeTags": ["license:apache-2.0", "library:lora", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:edge-inference", "custom_tag:prerouter", "custom_tag:lora", "custom_tag:ssd-offload"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 4597698253}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T12:59:29.226923+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4892870", "taskType": "text-generation", "verifyResult": null}
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "ManniX-ITA/Ornith-1.5-27B-A3B-CoderX", "modelProfile": {"architectures": ["Qwen3_5MoeForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 52429250296, "estimatedRequiredGiB": 58.623, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe_text", "modelscopeFileSize": 52454809630, "modelscopeLicense": "apache-2.0", "modelscopeParams": 26213016704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:expert-pruning", "custom_tag:code", "custom_tag:mtp", "custom_tag:ornith", "custom_tag:reap", "custom_tag:ream", "custom_tag:omnimergekit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 52454809833}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-16T13:37:27.004498+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4893449", "taskType": "text-generation", "verifyResult": null}