基于曦望官方 vLLM 镜像构建模型服务镜像

This commit is contained in:
luganyuan
2026-08-19 10:35:19 +08:00
parent e911988e4f
commit 5038ab1e36
4 changed files with 158 additions and 1 deletions

25
detect_tokenizer.py Normal file
View File

@@ -0,0 +1,25 @@
import os
import json
def detect(model_dir):
cfg_path = os.path.join(model_dir, "tokenizer_config.json")
if os.path.exists(cfg_path):
with open(cfg_path) as f:
cfg = json.load(f)
cls = cfg.get("tokenizer_class", "")
else:
cls = ""
files = os.listdir(model_dir)
if "tokenizer.json" in files:
return "fast", cls
if "tokenizer.model" in files:
return "sentencepiece", cls
if "vocab.json" in files and "merges.txt" in files:
return "bpe", cls
return "unknown", cls