Views
No views yet

| miniF2F-test | ProofNet | |
|---|---|---|
| ReProver | 26.5% | 13.8% |
| GPT-f | 36.6% | - |
| Hypertree Proof Search | 41.0% | - |
| InternLM2-StepProver | 54.5% | 18.1% |
| DeepSeek-Prover-V1 | 50.0% | - |
| DeepSeek-Prover-V1.5-Base | 42.2% | 13.2% |
| DeepSeek-Prover-V1.5-SFT | 57.4% | 22.9% |
| DeepSeek-Prover-V1.5-RL | 60.2% | 22.6% |
| DeepSeek-Prover-V1.5-RL + RMaxTS | 63.5% | 25.3% |
| Model | Download |
|---|---|
| DeepSeek-Prover-V1.5-Base | 🤗 HuggingFace |
| DeepSeek-Prover-V1.5-SFT | 🤗 HuggingFace |
| DeepSeek-Prover-V1.5-RL | 🤗 HuggingFace |
1@article{xin2024deepseekproverv15harnessingproofassistant,
2 title={DeepSeek-Prover-V1.5: Harnessing Proof Assistant Feedback for Reinforcement Learning and Monte-Carlo Tree Search},
3 author={Huajian Xin and Z. Z. Ren and Junxiao Song and Zhihong Shao and Wanjia Zhao and Haocheng Wang and Bo Liu and Liyue Zhang and Xuan Lu and Qiushi Du and Wenjun Gao and Qihao Zhu and Dejian Yang and Zhibin Gou and Z. F. Wu and Fuli Luo and Chong Ruan},
4 year={2024},
5 eprint={2408.08152},
6 archivePrefix={arXiv},
7 primaryClass={cs.CL},
8 url={https://arxiv.org/abs/2408.08152},
9}