Views
No views yet
| Training Loss | Epoch | Step | Validation Loss |
|---|---|---|---|
| 1.3883 | 0.9964 | 137 | 1.2905 |
| 1.0024 | 2.0 | 275 | 0.8735 |
| 0.4672 | 2.9964 | 412 | 0.6598 |
| 0.3044 | 4.0 | 550 | 0.5674 |
| 0.2501 | 4.9964 | 687 | 0.5263 |
| 0.5557 | 5.9782 | 822 | 0.5240 |
1@misc{BioMistral_fine_tuned,
2 title = {daphne604/{B}io{M}istral\_{D}{S}\_fine\_tuned · {H}ugging {F}ace --- huggingface.co},
3 author = {Daphne},
4 year = {2024},
5 publisher = {Hugging Face},
6 journal = {Hugging Face repository},
7 howpublished = {\url{https://huggingface.co/daphne604/BioMistral_DS_fine_tuned}}
8}
9
10@misc{vonwerra2022trl,
11 title = {{TRL: Transformer Reinforcement Learning}},
12 author = {Leandro von Werra and Younes Belkada and Lewis Tunstall and Edward Beeching and Tristan Thrush and Nathan Lambert and Shengyi Huang and Kashif Rasul and Quentin Gallouédec},
13 year = 2020,
14 journal = {GitHub repository},
15 publisher = {GitHub},
16 howpublished = {\url{https://github.com/huggingface/trl}}
17}