28 lines
1.1 KiB
JSON
28 lines
1.1 KiB
JSON
|
|
{
|
||
|
|
"repo_id": "SeongryongJung/Qwen3-4B-Material-GRPO-TR",
|
||
|
|
"output_dir": "/mnt/mole/SDPO/L2T/hf_upload_tr/Qwen3-4B-Material-GRPO-TR",
|
||
|
|
"experiment": "qwen3gen-material-GRPO-Qwen-Qwen3-4B-mbs8-train32-rollout8-lr1e-6-vllm0.8",
|
||
|
|
"best_step": 60,
|
||
|
|
"best_val_mean16": 0.7659574468085106,
|
||
|
|
"best_actor_dir": "/mnt/mole/SDPO/L2T/checkpoints/datasets/sciknoweval/material/qwen3gen-material-GRPO-Qwen-Qwen3-4B-mbs8-train32-rollout8-lr1e-6-vllm0.8/_actor_archive/global_step_60/actor",
|
||
|
|
"final_step": 100,
|
||
|
|
"final_val_mean16": 0.7626329787234043,
|
||
|
|
"final_actor_dir": "/mnt/mole/SDPO/L2T/checkpoints/datasets/sciknoweval/material/qwen3gen-material-GRPO-Qwen-Qwen3-4B-mbs8-train32-rollout8-lr1e-6-vllm0.8/global_step_100/actor",
|
||
|
|
"train_rows": 90,
|
||
|
|
"val_rows": 10,
|
||
|
|
"hf_model_files": [
|
||
|
|
"added_tokens.json",
|
||
|
|
"chat_template.jinja",
|
||
|
|
"config.json",
|
||
|
|
"generation_config.json",
|
||
|
|
"merges.txt",
|
||
|
|
"model.safetensors.index.json",
|
||
|
|
"special_tokens_map.json",
|
||
|
|
"tokenizer.json",
|
||
|
|
"tokenizer_config.json",
|
||
|
|
"vocab.json",
|
||
|
|
"model-00001-of-00002.safetensors",
|
||
|
|
"model-00002-of-00002.safetensors"
|
||
|
|
]
|
||
|
|
}
|