@misc{indiciae94b2c2f7adb7, title = {Legal Mathematical Reasoning with LLMs: Procedural Alignment through Two-Stage Reinforcement Learning}, author = {Kepu Zhang and Guofu Xie and Weijie Yu and Mingyue Xu and Xu Tang and Yaxin Li and Jun Xu}, year = {2025}, url = {https://arxiv.org/abs/2504.02590}, note = {Source identifier: 2504.02590} }