Views
No views yet
1@misc{li2024calibratingllmspreferenceoptimization,
2 title={Calibrating LLMs with Preference Optimization on Thought Trees for Generating Rationale in Science Question Scoring},
3 author={Jiazheng Li and Hainiu Xu and Zhaoyue Sun and Yuxiang Zhou and David West and Cesare Aloisi and Yulan He},
4 year={2024},
5 eprint={2406.19949},
6 archivePrefix={arXiv},
7 primaryClass={cs.CL},
8 url={https://arxiv.org/abs/2406.19949},
9}| Training Loss | Epoch | Step | Validation Loss |
|---|---|---|---|
| 0.9813 | 0.63 | 100 | 0.9671 |
| 0.9108 | 1.26 | 200 | 0.9250 |
| 0.8976 | 1.9 | 300 | 0.9091 |
| 0.8687 | 2.53 | 400 | 0.9005 |
| 0.8548 | 3.16 | 500 | 0.8958 |
| 0.8468 | 3.79 | 600 | 0.8952 |