@misc{indiciaea504011122c8, title = {A Practical Two-Stage Recipe for Mathematical LLMs: Maximizing Accuracy with SFT and Efficiency with Reinforcement Learning}, author = {Hiroshi Yoshihara and Taiki Yamaguchi and Yuichi Inoue}, year = {2025}, url = {https://arxiv.org/abs/2507.08267}, note = {Source identifier: 2507.08267} }