初始化项目,由ModelHub XC社区提供模型

Model: wvnvwn/llama-2-13b-chat-hf-gsm8k-sn-tuned-lr5e-5
Source: Original Platform
This commit is contained in:
ModelHub XC
2026-08-09 01:07:18 +08:00
commit c959c1de52
18 changed files with 277768 additions and 0 deletions

21
finetune_config.json Normal file
View File

@@ -0,0 +1,21 @@
{
"base_model": "wvnvwn/llama-2-13b-chat-hf-only-sn-tuned-lr5e-5",
"fine_tuning_type": "GSM8K Fine-tuning with Safety Neuron Freezing",
"safety_neurons_file": "/NHNHOME/WORKSPACE/26msit001_A/edge_ai_lab/jongbokwon/minseong_results/safety_llama-2-13b-chat-hf/safety_neuron.txt",
"dataset": "GSM8K",
"num_train_samples": 7473,
"batch_size": 4,
"grad_accum": 4,
"learning_rate": 5e-05,
"weight_decay": 0.01,
"warmup_ratio": 0.1,
"epochs": 3,
"max_length": 1024,
"max_grad_norm": 1.0,
"lr_scheduler_type": "cosine",
"optimizer": "adamw_torch",
"gradient_checkpointing": false,
"dtype": "bf16",
"trainer_type": "Trainer",
"strategy": "Freeze safety neurons, train others"
}