state: generation 8676 (result)
This commit is contained in:
@@ -600,3 +600,4 @@
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "rapid-mlx/Qwen3.8-27B-4bit-MTP-MLX", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16293475486, "estimatedRequiredGiB": 18.239, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 16320427112, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4665462000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:rapid-mlx", "custom_tag:qwen3.8", "custom_tag:qwen3.8-27b", "custom_tag:4-bit", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16320427112}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T10:12:07.489740+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4913272", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "ManniX-ITA/Ornith-1.5-27B-A3B-CoderX", "modelProfile": {"architectures": ["Qwen3_5MoeForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 52429250296, "estimatedRequiredGiB": 58.623, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5_moe_text", "modelscopeFileSize": 52454809833, "modelscopeLicense": "apache-2.0", "modelscopeParams": 26213016704, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5_moe_text", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:moe", "custom_tag:expert-pruning", "custom_tag:code", "custom_tag:mtp", "custom_tag:ornith", "custom_tag:reap", "custom_tag:ream", "custom_tag:omnimergekit"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 52454809833}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T10:12:07.491666+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4913273", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": null, "modelId": "ornith-ai/Ornith-1.5-9B-MLX-6bit", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7276346365, "estimatedRequiredGiB": 8.156, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 7298158256, "modelscopeLicense": null, "modelscopeParams": 1959473664, "modelscopeTags": ["model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 7298158256}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T10:12:07.487236+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4913275", "taskType": "text-generation", "verifyResult": null}
|
||||
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "RentedNoodle/Qwen3.8-27B-OrcaRouter-GSQ-RCO-IQ3_XXS-GGUF", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "289bdbc89807736c7fc9eba7819ba7c6f178bbf5fb1f56a5683ad9b7fa830bfd", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 629247648, "estimatedRequiredGiB": 25.142, "gpuMemoryEvidence": {"memoryGiB": 48.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 48.0, "maximumRepositorySizeGiB": 40.0, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 22496243465, "modelscopeLicense": "apache-2.0", "modelscopeParams": 27320697856, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:gguf", "task:text-generation", "custom_tag:qwen3", "custom_tag:qwen3.8", "custom_tag:orcarouter", "custom_tag:gsq", "custom_tag:rco", "custom_tag:gguf", "custom_tag:quantized", "custom_tag:speculative-decoding", "custom_tag:mtp", "custom_tag:text-generation", "custom_tag:long-context", "custom_tag:tool-use", "custom_tag:16gb-vram", "custom_tag:image-text-to-text", "custom_tag:vision", "custom_tag:multimodal", "custom_tag:mmproj", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 22496243465}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-17T10:17:26.606251+00:00", "targetGpu": "Mthreads_s4000", "taskId": "4913350", "taskType": "text-generation", "verifyResult": null}
|
||||
|
||||
Reference in New Issue
Block a user