-
Concise reasoning: 65% fewer output tokens across five math benchmarks, with full reasoning preserved.
-
Performance: 15% higher accuracy than alternative solutions.
-
Efficiency: 20% faster compared with the GRPO training.
1@article{gao2025concise,
2 title={Concise Reasoning in the Lens of Lagrangian Optimization},
3 author={Gao, Chengqian and Li, Haonan and Killian, Taylor W and She, Jianshu and Wang, Renxi and Ma, Liqun and Cheng, Zhoujun and Hao, Shibo and Xu, Zhiqiang},
4 journal={arXiv preprint arXiv:2510.10168},
5 year={2025}
6}