diff --git a/.modelhub_state/recovery_intents.jsonl b/.modelhub_state/recovery_intents.jsonl index 13dc9de4..b584e61e 100644 --- a/.modelhub_state/recovery_intents.jsonl +++ b/.modelhub_state/recovery_intents.jsonl @@ -673,7 +673,7 @@ {"batchId": "3cfa37416a9349629e882947d2816a4f", "completedAt": "2026-09-11T10:27:40.997258+00:00", "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configSource": "modelhub_live", "createdAt": "2026-09-11T10:27:32.517702+00:00", "framework": "vllm_0_17_0_corex_4_4_0", "intentId": "faf0e0acf87c4526ad73c64cde85600e", "lastModified": "2026-09-09T16:53:28+00:00", "modelAddress": "https://modelscope.cn/models/laion/swesmith-nl2bash-stack-bugsseq", "reason": null, "reconciledAt": "2026-09-11T15:36:01.451074+00:00", "repoId": "laion/swesmith-nl2bash-stack-bugsseq", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Iluvatar_bi-150", "taskId": "4776575", "taskType": "text-generation"} {"batchId": "3cfa37416a9349629e882947d2816a4f", "completedAt": "2026-09-11T10:27:40.997261+00:00", "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configSource": "modelhub_live", "createdAt": "2026-09-11T10:27:32.517746+00:00", "framework": "vllm_0_17_0_corex_4_4_0", "intentId": "230e7c11b34e45f8ba042c7e6117d732", "lastModified": "2026-08-11T14:42:47+00:00", "modelAddress": "https://modelscope.cn/models/BAAI/AREX-Turbo", "reason": null, "reconciledAt": "2026-09-11T15:36:01.451503+00:00", "repoId": "BAAI/AREX-Turbo", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "recovered_active", "targetGpu": "Iluvatar_bi-150", "taskId": "4776578", "taskType": "text-generation"} {"batchId": "cf90bf18e1d046ddadc8968a623bf967", "completedAt": "2026-09-11T15:31:13.164650+00:00", "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configSource": "modelhub_live", "createdAt": "2026-09-11T15:30:57.117504+00:00", "framework": "vllm_tokenizer_patch", "intentId": "6a206bff22bf42d298dbae81668fdd2c", "lastModified": "2026-09-09T06:50:48+00:00", "modelAddress": "https://modelscope.cn/models/ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "reason": null, "reconciledAt": "2026-09-11T15:36:01.451500+00:00", "repoId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "safeConfigVector": {"gpuNum": 1}, "status": "recovered_active", "targetGpu": "Ascend_910-b4", "taskId": "4782245", "taskType": "text-generation"} -{"batchId": "d5d8a922f1b942988202e0fdc139eda4", "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configSource": "modelhub_live", "createdAt": "2026-09-11T15:45:33.425511+00:00", "framework": "vllm_tokenizer_patch", "intentId": "190b274e1b10409e8152fff5833030f0", "lastModified": "2026-09-11T15:37:46+00:00", "modelAddress": "https://modelscope.cn/models/CohereLabs/tiny-aya-base-32K", "repoId": "CohereLabs/tiny-aya-base-32K", "safeConfigVector": {"gpuNum": 1}, "status": "pending", "targetGpu": "Ascend_910-b4", "taskType": "text-generation"} +{"batchId": "d5d8a922f1b942988202e0fdc139eda4", "completedAt": "2026-09-11T15:45:41.271246+00:00", "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configSource": "modelhub_live", "createdAt": "2026-09-11T15:45:33.425511+00:00", "framework": "vllm_tokenizer_patch", "intentId": "190b274e1b10409e8152fff5833030f0", "lastModified": "2026-09-11T15:37:46+00:00", "modelAddress": "https://modelscope.cn/models/CohereLabs/tiny-aya-base-32K", "reason": null, "repoId": "CohereLabs/tiny-aya-base-32K", "safeConfigVector": {"gpuNum": 1}, "status": "submitted", "targetGpu": "Ascend_910-b4", "taskId": "4782449", "taskType": "text-generation"} {"batchId": "3cfa37416a9349629e882947d2816a4f", "completedAt": "2026-09-11T10:27:40.997245+00:00", "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configSource": "modelhub_live", "createdAt": "2026-09-11T10:27:32.517524+00:00", "framework": "vllm_0_17_0_corex_4_4_0", "intentId": "11684395861643dd8cc4c1610b8518a9", "repoId": "nv-community/Qwen3.8-27B-NVFP4", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "Iluvatar_bi-150", "taskId": null, "taskType": "text-generation"} {"batchId": "3cfa37416a9349629e882947d2816a4f", "completedAt": "2026-09-11T10:27:40.997200+00:00", "configFingerprint": "e57ef22c2d8b2b75e495765c383ba0c5b795fe18d7b7f1baf8dcfadd87065bd8", "configSource": "modelhub_live", "createdAt": "2026-09-11T10:27:32.516990+00:00", "framework": "llamacpp", "intentId": "a0656dc70ed9494580076f5910d1e348", "repoId": "TokenRhythm/NeoHorse-1-4B-GGUF", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "Iluvatar_bi-150", "taskId": null, "taskType": "text-generation"} {"batchId": "75d7acb27d704801b2ce084de0bf3900", "completedAt": "2026-09-11T02:14:19.417302+00:00", "configFingerprint": "3487754ba6ce87b4107d75c6476a4c5e5528f9179ac18665ecda0531c247da5e", "configSource": "modelhub_live", "createdAt": "2026-09-11T02:14:10.788042+00:00", "framework": "llamacpp", "intentId": "a31ea00fe3a847e6b52279c71b1fd09a", "repoId": "TokenRhythm/NeoHorse-1-4B-GGUF", "safeConfigVector": {"gpuNum": 1, "maxModelLen": 4096}, "status": "uniqueness_rejected", "targetGpu": "Ascend_910-b4", "taskId": null, "taskType": "text-generation"} diff --git a/.modelhub_state/routing_intelligence.json b/.modelhub_state/routing_intelligence.json index 2e54fe35..437de910 100644 --- a/.modelhub_state/routing_intelligence.json +++ b/.modelhub_state/routing_intelligence.json @@ -1,13 +1,13 @@ { "acceptedByCategory": { - "unified_success_first": 808 + "unified_success_first": 809 }, "acceptedByRoute": { "text-generation|Ascend_910-b3|llamacpp": 11, "text-generation|Ascend_910-b3|vllm": 16, "text-generation|Ascend_910-b3|vllm_tokenizer_patch": 20, "text-generation|Ascend_910-b4|llamacpp": 8, - "text-generation|Ascend_910-b4|vllm_tokenizer_patch": 31, + "text-generation|Ascend_910-b4|vllm_tokenizer_patch": 32, "text-generation|Biren_166m|vllm": 52, "text-generation|Cambricon_mlu-370-x4|vllm-mlu": 28, "text-generation|Cambricon_mlu-370-x8|vllm": 13, @@ -26,8 +26,8 @@ "text-generation|hygon_k100-ai|llamacpp": 11, "text-generation|hygon_k100-ai|vllm-patch-tokenizer": 45 }, - "acceptedSinceRefresh": 808, - "acceptedTotal": 808, - "generatedAt": "2026-09-11T15:31:13.099018+00:00", + "acceptedSinceRefresh": 809, + "acceptedTotal": 809, + "generatedAt": "2026-09-11T15:45:41.242557+00:00", "version": 1 } diff --git a/ledger/submissions.jsonl b/ledger/submissions.jsonl index b1931686..71e6b797 100644 --- a/ledger/submissions.jsonl +++ b/ledger/submissions.jsonl @@ -765,3 +765,4 @@ {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/JANGQ-AI/AppleScript-8B-JANG_4M", "modelId": "JANGQ-AI/AppleScript-8B-JANG_4M", "submitTime": "2026-09-08T15:58:20.200513+00:00", "targetGpu": "Vastai_va16", "taskId": "4722293", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-vastai-va16"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/OpenBMB/BitCPM-CANN-8B", "modelId": "OpenBMB/BitCPM-CANN-8B", "submitTime": "2026-09-09T12:30:12.314907+00:00", "targetGpu": "Vastai_va16", "taskId": "4740401", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-vastai-va16"} {"framework": "vllm_fix_tokenizer", "modelAddress": "https://modelscope.cn/models/Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16", "modelId": "Youssofal/Qwen3.8-27B-MTPLX-Bare-Speed-FP16", "submitTime": "2026-09-09T21:08:22.097030+00:00", "targetGpu": "Kunlunxin_p-800", "taskId": "4746441", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-fix-tokenizer-kunlunxin-p-800"} +{"framework": "vllm_tokenizer_patch", "modelAddress": "https://modelscope.cn/models/CohereLabs/tiny-aya-base-32K", "modelId": "CohereLabs/tiny-aya-base-32K", "submitTime": "2026-09-11T15:45:34.084559+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4782449", "taskType": "text-generation", "templateId": "modelhub-live-text-generation-vllm-tokenizer-patch-ascend-910-b4"} diff --git a/manifest.json b/manifest.json index 21ac262a..3157ca24 100644 --- a/manifest.json +++ b/manifest.json @@ -9,16 +9,16 @@ ".modelhub_state/queue_cleanup_latest.json": "aef9d56345e6e6b2c6e9bcc5c224abe9dc0699f57e994e20e663ac3e21dd4b56", ".modelhub_state/recent_outcomes.jsonl": "120eccf9fa38b19d315a7cd6d5484fed6a29bc5c3a6301a857e192af6e539172", ".modelhub_state/recovery_active_tasks.jsonl": "395379a87de5adea66ce3a6b27e840315fd427a8daab264f391eb461e6220ed3", - ".modelhub_state/recovery_intents.jsonl": "ac12f1d8d00680219a78f920da126dae375b632d8e23f326b40a69de8063ba12", - ".modelhub_state/routing_intelligence.json": "42c5411972287c8742c3908f0f7c8569ff4293a615cd0886bb6bfdb0fd57171d", + ".modelhub_state/recovery_intents.jsonl": "4023021294a90b19f5bcea7b49eb9f4d99d52b5e0188b31ac071f21b05edd0ec", + ".modelhub_state/routing_intelligence.json": "c1b5e83c44e289b6950ffb5b9fab479b4a9b0e65b093861ea00f4e6d64a86a81", ".modelhub_state/submission_exclusions.jsonl": "3790b5c148f222a4de26d50a3d8afc58891c8da251a82d20dbaf4fc35e6dbd89", ".modelhub_state/worker_crashes.jsonl": "7fa483580e87374493226e80bcf40bf8044847acba70045c96fa175966a1a8ba", - "ledger/submissions.jsonl": "d437711128f3914d9360a660940d87da2fc558cd5f7343f697844791c650f6b5", - "outcomes/submissions.jsonl": "03e65b2a362b7ad47fa25b1d15b548c3546346f4698694edd9269267b1a646c9" + "ledger/submissions.jsonl": "eb74fa4ea8a8a6c1682fff299594d7e16e42c0ec62596fdc585b96affefa001e", + "outcomes/submissions.jsonl": "8a709965c2e3eabdb597486d8ef12a0931762f64c48ea3b7a011d8eff4791cec" }, - "generation": 5752, - "phase": "intent", + "generation": 5753, + "phase": "result", "schemaVersion": 1, - "updatedAt": "2026-09-11T15:45:33.486792+00:00", + "updatedAt": "2026-09-11T15:45:41.361882+00:00", "writerId": "18252d8a8ef94333abd55823b2d81c42" } diff --git a/outcomes/submissions.jsonl b/outcomes/submissions.jsonl index 8b6bef37..05273f03 100644 --- a/outcomes/submissions.jsonl +++ b/outcomes/submissions.jsonl @@ -605,3 +605,4 @@ {"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": null, "modelId": "laion/swesmith-nl2bash-stack-bugsseq", "modelProfile": {"architectures": ["Qwen3ForCausalLM"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 8593, "estimatedRequiredGiB": 18.327, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3", "modelscopeFileSize": 16398873171, "modelscopeLicense": "apache-2.0", "modelscopeParams": 8190735360, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:llama-factory", "custom_tag:full", "custom_tag:generated_from_trainer"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 16398873171}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T10:27:40.674298+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4776575", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_0_17_0_corex_4_4_0", "lastSyncTime": null, "modelId": "BAAI/AREX-Turbo", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "1cc16a7df03cf7b1485e2fda7e5c05e286a546d04e80bdd595c015004b3852a8", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 9078620368, "estimatedRequiredGiB": 10.18, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "modelhub_preflight_oom:24"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 9109259594, "modelscopeLicense": "apache-2.0", "modelscopeParams": 4539265536, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:agent", "custom_tag:deep-research", "custom_tag:reasoning", "custom_tag:tool-use", "custom_tag:long-context", "custom_tag:qwen3.5", "custom_tag:dense"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 9109259594}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T10:27:40.670287+00:00", "targetGpu": "Iluvatar_bi-150", "taskId": "4776578", "taskType": "text-generation", "verifyResult": null} {"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "ddalcu/Qwen3.8-27B-MLX-Serve-iQ-MLX-3.8bpw", "modelProfile": {"architectures": ["Qwen3_5ForConditionalGeneration"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 12999919368, "estimatedRequiredGiB": 14.555, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "qwen3_5", "modelscopeFileSize": 13023155398, "modelscopeLicense": "apache-2.0", "modelscopeParams": 3544073216, "modelscopeTags": ["license:apache-2.0", "model_type:qwen3_5", "library:mlx", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:mlx", "custom_tag:mlx-serve", "custom_tag:qwen3_5", "custom_tag:apple-silicon", "custom_tag:imatrix", "custom_tag:mtp", "custom_tag:speculative-decoding"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 13023155398}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T15:30:58.007700+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4782245", "taskType": "text-generation", "verifyResult": null} +{"failReason": null, "framework": "vllm_tokenizer_patch", "lastSyncTime": null, "modelId": "CohereLabs/tiny-aya-base-32K", "modelProfile": {"architectures": ["Cohere2ForCausalLM"], "configFingerprint": "167e80c9a74f5acd91819ac9a661546c8992cc5beb0bdd0a1dc9fc9659a5e778", "configOptimization": {"applied": false, "source": "official"}, "configSource": "modelhub_live", "estimatedLoadBytes": 6700539432, "estimatedRequiredGiB": 7.522, "gpuMemoryEvidence": {"memoryGiB": 32.0, "source": "local_modelhub_preflight_oom"}, "gpuMemoryGiB": 32.0, "maximumRepositorySizeGiB": 26.667, "memorySizingBasis": "recursive_repository_on_disk", "modelCard": {}, "modelType": "cohere2", "modelscopeFileSize": 3443, "modelscopeLicense": "cc-by-nc-4.0", "modelscopeParams": 3350243328, "modelscopeTags": ["license:cc-by-nc-4.0", "model_type:cohere2", "library:transformer", "library:safetensors", "library:pytorch", "task:text-generation", "custom_tag:aya", "custom_tag:multilingual", "custom_tag:long-context", "custom_tag:tiny-aya"], "modelscopeTasks": ["text-generation"], "quantizationMethod": null, "repositoryOnDiskBytes": 6730839322}, "outcome": "pending", "status": "pending", "submitTime": "2026-09-11T15:45:34.084559+00:00", "targetGpu": "Ascend_910-b4", "taskId": "4782449", "taskType": "text-generation", "verifyResult": null}