Files
zerorlvrcode-qwen2.5-1.5b/eval-results/ifbench/metrics.json

16 lines
430 B
JSON
Raw Normal View History

{
"ifbench": {
"pass@1": {
"num_prompts": 294,
"num_instructions": 335,
"average_score": 20.91912884556808,
"prompt_strict_accuracy": 18.70748299319728,
"instruction_strict_accuracy": 20.8955223880597,
"prompt_loose_accuracy": 21.08843537414966,
"instruction_loose_accuracy": 22.98507462686567,
"num_entries": 294,
"avg_tokens": 1612,
"gen_seconds": 37
}
}
}