初始化项目,由ModelHub XC社区提供模型
Model: thu-ml/STAIR-Qwen2-7B-DPO-3 Source: Original Platform
This commit is contained in:
4
trainer_log.jsonl
Normal file
4
trainer_log.jsonl
Normal file
@@ -0,0 +1,4 @@
|
||||
{"current_steps": 10, "total_steps": 31, "loss": 0.8699, "accuracy": 0.48750001192092896, "learning_rate": 4.415111107797445e-07, "epoch": 0.32, "percentage": 32.26, "elapsed_time": "0:02:41", "remaining_time": "0:05:38"}
|
||||
{"current_steps": 20, "total_steps": 31, "loss": 0.7936, "accuracy": 0.699999988079071, "learning_rate": 1.782991918222275e-07, "epoch": 0.64, "percentage": 64.52, "elapsed_time": "0:05:21", "remaining_time": "0:02:57"}
|
||||
{"current_steps": 30, "total_steps": 31, "loss": 0.7511, "accuracy": 0.668749988079071, "learning_rate": 1.690410564514244e-09, "epoch": 0.96, "percentage": 96.77, "elapsed_time": "0:08:09", "remaining_time": "0:00:16"}
|
||||
{"current_steps": 31, "total_steps": 31, "epoch": 0.992, "percentage": 100.0, "elapsed_time": "0:08:49", "remaining_time": "0:00:00"}
|
||||
Reference in New Issue
Block a user