|
|
|
|
@@ -289,7 +289,6 @@
|
|
|
|
|
{"failReason": null, "framework": "vllm-patch-tokenizer", "lastSyncTime": "2026-09-20T13:18:18.977292+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "0d5bc82787a89d4313672808d1a8629a1fc675c91aa6a08865f770aa24447fbc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T05:16:13.324611+00:00", "targetGpu": "hygon_k100-ai", "taskId": "4971311", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T14:18:03.258951+00:00", "modelId": "RedHatAI/gemma-2-27b-it-quantized.w8a16", "modelProfile": {"architectures": ["Gemma2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 30775446656, "estimatedRequiredGiB": 34.419, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma2", "modelscopeFileSize": 30797371689, "modelscopeLicense": "gemma", "modelscopeParams": 28406776320, "modelscopeTags": ["license:gemma", "model_type:gemma2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:int8", "custom_tag:vllm"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 30797371689}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T06:16:07.269550+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4972172", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T16:51:24.761939+00:00", "modelId": "inceptionai/Jais-2-8B-Chat", "modelProfile": {"architectures": ["Jais2ForCausalLM"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 16180861064, "estimatedRequiredGiB": 18.097, "gpuMemoryEvidence": {"memoryGiB": 100.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 100.0, "maximumRepositorySizeGiB": 83.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "jais2", "modelscopeFileSize": 16193187697, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8090401280, "modelscopeTags": ["license:apache-2.0", "model_type:jais2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16193187697}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T08:49:13.744928+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4974217", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm", "lastSyncTime": "2026-09-20T17:54:13.170240+00:00", "modelId": "bharatgenai/Param2-17B-A2.4B-Thinking", "modelProfile": {"architectures": ["Param2MoEForCausalLM"], "configFingerprint": "53aadead5bb69dd6c72976cf6353832080116244f86ffae86297646e5bb48fdc", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 34302758832, "estimatedRequiredGiB": 38.353, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "param2moe", "modelscopeFileSize": 34317809830, "modelscopeLicense": null, "modelscopeParams": 17151125376, "modelscopeTags": ["model_type:param2moe", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:mixture-of-experts", "custom_tag:multilingual", "custom_tag:indian-languages"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 34317809830}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T09:41:34.108703+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4976623", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_fix_tokenizer", "lastSyncTime": "2026-09-20T21:06:45.850274+00:00", "modelId": "RedHatAI/gemma-4-12B-it-FP8-Dynamic", "modelProfile": {"architectures": ["Gemma4UnifiedForConditionalGeneration"], "configFingerprint": "16f74ea0f409861ab0c8169a25d9484cb482e27d6a3ad7c529a70b7cdbfabee3", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 15037393456, "estimatedRequiredGiB": 16.842, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "gemma4_unified", "modelscopeFileSize": 15069601989, "modelscopeLicense": "apache-2.0", "modelscopeParams": 12966363184, "modelscopeTags": ["license:apache-2.0", "model_type:gemma4_unified", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "task:image-text-to-text", "custom_tag:fp8", "custom_tag:vllm", "custom_tag:llm-compressor", "custom_tag:compressed-tensors"], "modelscopeTasks": ["text-generation", "image-text-to-text"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 15069601989}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T13:01:42.205703+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4979280", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T22:43:47.754909+00:00", "modelId": "neuralmagic/starcoder2-3b-quantized.w8a8", "modelProfile": {"architectures": ["Starcoder2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 4093158536, "estimatedRequiredGiB": 4.578, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "starcoder2", "modelscopeFileSize": 4096457730, "modelscopeLicense": "bigcode-openrail-m", "modelscopeParams": 3181366272, "modelscopeTags": ["license:bigcode-openrail-m", "model_type:starcoder2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:code"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 4096457730}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T14:38:07.337494+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4980575", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-09-20T23:08:35.573457+00:00", "modelId": "neuralmagic/Qwen2-1.5B-Instruct-quantized.w8a8", "modelProfile": {"architectures": ["Qwen2ForCausalLM"], "configFingerprint": "7b5500238d52e9e2a21471d722bc1636ca047e1a12ba49c491593b86bec602b8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2245331488, "estimatedRequiredGiB": 2.522, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen2", "modelscopeFileSize": 2256941144, "modelscopeLicense": "apache-2.0", "modelscopeParams": 1777088000, "modelscopeTags": ["license:apache-2.0", "model_type:qwen2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation"], "modelscopeTasks": ["text-generation"], "quantizationMethod": "compressed-tensors", "repositoryOnDiskBytes": 2256941144}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-20T15:06:16.550878+00:00", "targetGpu": "Ascend_910-b3", "taskId": "4980890", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
@@ -749,9 +748,9 @@
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T02:54:18.466070+00:00", "modelId": "webAI-Official/TwIL-LM3-Pro", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 7319517272, "estimatedRequiredGiB": 29.512, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:granite-4.2", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 26407018955}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T18:43:37.348780+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200237", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T02:54:18.466091+00:00", "modelId": "webAI-Official/TwIL-LM2", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5067121544, "estimatedRequiredGiB": 20.413, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:twil-lm", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18264867356}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T18:43:37.346955+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200239", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T02:54:18.466100+00:00", "modelId": "RWKV/RWKV7-G1k-2.9B-20260930", "modelProfile": {"architectures": ["Rwkv7ForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5896236664, "estimatedRequiredGiB": 6.592, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "rwkv7", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:rwkv7", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:rwkv", "custom_tag:rwkv7", "custom_tag:recurrent", "custom_tag:causal-lm", "custom_tag:conversational"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5898355154}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T18:43:37.247818+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200236", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "mlx-community/K2-Horizon-3.7B-8bit", "modelProfile": {"architectures": ["K2HorizonForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5374667178, "estimatedRequiredGiB": 6.03, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "k2_horizon", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:k2_horizon", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:dense", "custom_tag:k2-horizon"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5395463319}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:16.876280+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200515", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "webAI-Official/TwIL-LM3-Pro", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "645727a8c34425a2c3e191b5407d247f231e0d4e88f0bd3deb19333f39d09ba4", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3892651456, "estimatedRequiredGiB": 29.512, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:granite-4.2", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 26407018955}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.049878+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200516", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "llamacpp", "lastSyncTime": null, "modelId": "webAI-Official/TwIL-LM2", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8120661075bb80389dd6d7b11ccbb58a0586dc072989d5570471de11994360d9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2694119616, "estimatedRequiredGiB": 20.413, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:twil-lm", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18264867356}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.051413+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200518", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "bench-labs/pulvis-v1", "modelProfile": {"architectures": ["PulvisForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 11853976, "estimatedRequiredGiB": 0.123, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "pulvis", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8886720, "modelscopeTags": ["license:apache-2.0", "model_type:pulvis", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:causal-lm", "custom_tag:language-model", "custom_tag:base-model", "custom_tag:pretrained-from-scratch", "custom_tag:small-language-model"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 109820987}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.054397+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200517", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "RWKV/RWKV7-G1k-2.9B-20260930", "modelProfile": {"architectures": ["Rwkv7ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5896236664, "estimatedRequiredGiB": 6.592, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "rwkv7", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:rwkv7", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:rwkv", "custom_tag:rwkv7", "custom_tag:recurrent", "custom_tag:causal-lm", "custom_tag:conversational"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5898355154}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:00:30.056861+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200519", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T03:00:54.465099+00:00", "modelId": "mlx-community/K2-Horizon-3.7B-8bit", "modelProfile": {"architectures": ["K2HorizonForCausalLM"], "configFingerprint": "8be5b78bdba5c4413c1e4ab91a46e27c34642b25d5dc059ad3a11d50d7e33593", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5374667178, "estimatedRequiredGiB": 6.03, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "k2_horizon", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:k2_horizon", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:dense", "custom_tag:k2-horizon"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5395463319}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:16.876280+00:00", "targetGpu": "MetaX_c-500", "taskId": "5200515", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-10-01T03:00:54.465090+00:00", "modelId": "webAI-Official/TwIL-LM3-Pro", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "645727a8c34425a2c3e191b5407d247f231e0d4e88f0bd3deb19333f39d09ba4", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 3892651456, "estimatedRequiredGiB": 29.512, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:granite-4.2", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 26407018955}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.049878+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200516", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "llamacpp", "lastSyncTime": "2026-10-01T03:00:54.465128+00:00", "modelId": "webAI-Official/TwIL-LM2", "modelProfile": {"architectures": ["GraniteForCausalLM"], "configFingerprint": "8120661075bb80389dd6d7b11ccbb58a0586dc072989d5570471de11994360d9", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 2694119616, "estimatedRequiredGiB": 20.413, "gpuMemoryEvidence": {"memoryGiB": 64.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 64.0, "maximumRepositorySizeGiB": 53.333, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "granite", "modelscopeFileSize": null, "modelscopeLicense": "other", "modelscopeParams": null, "modelscopeTags": ["license:other", "model_type:granite", "library:pytorch", "library:transformer", "library:lora", "library:safetensors", "library:gguf", "task:text-generation", "custom_tag:granite", "custom_tag:formal-logic", "custom_tag:reasoning", "custom_tag:lora", "custom_tag:model-merging", "custom_tag:wise-ft", "custom_tag:reinforcement-learning", "custom_tag:grpo", "custom_tag:twil-lm", "custom_tag:gguf", "deploy:swingdeploy"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 18264867356}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.051413+00:00", "targetGpu": "hygon_k100-ai", "taskId": "5200518", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T03:00:54.465107+00:00", "modelId": "bench-labs/pulvis-v1", "modelProfile": {"architectures": ["PulvisForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 11853976, "estimatedRequiredGiB": 0.123, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "pulvis", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8886720, "modelscopeTags": ["license:apache-2.0", "model_type:pulvis", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:causal-lm", "custom_tag:language-model", "custom_tag:base-model", "custom_tag:pretrained-from-scratch", "custom_tag:small-language-model"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 109820987}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.054397+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200517", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": "2026-10-01T03:00:54.465067+00:00", "modelId": "RWKV/RWKV7-G1k-2.9B-20260930", "modelProfile": {"architectures": ["Rwkv7ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5896236664, "estimatedRequiredGiB": 6.592, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "rwkv7", "modelscopeFileSize": null, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:rwkv7", "library:pytorch", "library:transformer", "library:safetensors", "task:text-generation", "custom_tag:rwkv", "custom_tag:rwkv7", "custom_tag:recurrent", "custom_tag:causal-lm", "custom_tag:conversational"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5898355154}, "outcome": "pending", "status": "waiting", "submitTime": "2026-09-30T19:00:30.056861+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200519", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "mlx-community/K2-Horizon-3.7B-8bit", "modelProfile": {"architectures": ["K2HorizonForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 5374667178, "estimatedRequiredGiB": 6.03, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "k2_horizon", "modelscopeFileSize": 5395463319, "modelscopeLicense": "apache-2.0", "modelscopeParams": null, "modelscopeTags": ["license:apache-2.0", "model_type:k2_horizon", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:quantized", "custom_tag:dense", "custom_tag:k2-horizon"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 5395463319}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-30T19:16:58.857252+00:00", "targetGpu": "Ascend_910-b4", "taskId": "5200707", "taskType": "text-generation", "verifyResult": null}
|
|
|
|
|
|