初始化项目,由ModelHub XC社区提供模型
Model: cortexso/qwen2 Source: Original Platform
This commit is contained in:
17
model.yml
Normal file
17
model.yml
Normal file
@@ -0,0 +1,17 @@
|
||||
name: qwen2
|
||||
model: qwen2:7B
|
||||
version: 1
|
||||
|
||||
# Results Preferences
|
||||
top_p: 0.95
|
||||
temperature: 0.7
|
||||
frequency_penalty: 0
|
||||
presence_penalty: 0
|
||||
max_tokens: 4096 # Infer from base config.json -> max_position_embeddings
|
||||
stream: true # true | false
|
||||
|
||||
# Engine / Model Settings
|
||||
ngl: 33 # Infer from base config.json -> num_attention_heads
|
||||
ctx_len: 4096 # Infer from base config.json -> max_position_embeddings
|
||||
engine: llama-cpp
|
||||
prompt_template: "<|im_start|>system\n{system_message}<|im_end|>\n<|im_start|>user\n{prompt}<|im_end|>\n<|im_start|>assistant"
|
||||
Reference in New Issue
Block a user