初始化项目,由ModelHub XC社区提供模型
Model: geoffmunn/Qwen3-8B-f16 Source: Original Platform
This commit is contained in:
25
MODELFILE
Normal file
25
MODELFILE
Normal file
@@ -0,0 +1,25 @@
|
||||
# MODELFILE for Qwen3-8B-GGUF
|
||||
# Used by LM Studio, OpenWebUI, GPT4All, etc.
|
||||
|
||||
context_length: 32768
|
||||
embedding: false
|
||||
f16: cpu
|
||||
|
||||
# Chat template using ChatML (used by Qwen)
|
||||
prompt_template: >-
|
||||
<|im_start|>system
|
||||
You are a helpful assistant.<|im_end|>
|
||||
<|im_start|>user
|
||||
{prompt}<|im_end|>
|
||||
<|im_start|>assistant
|
||||
|
||||
# Stop sequences help end generation cleanly
|
||||
stop: "<|im_end|>"
|
||||
stop: "<|im_start|>"
|
||||
|
||||
# Default sampling (optimized for thinking mode)
|
||||
temperature: 0.6
|
||||
top_p: 0.95
|
||||
top_k: 20
|
||||
min_p: 0.0
|
||||
repeat_penalty: 1.1
|
||||
Reference in New Issue
Block a user