@misc{indiciae69788f3cf1e3, title = {GRPO-LEAD: A Difficulty-Aware Reinforcement Learning Approach for Concise Mathematical Reasoning in Language Models}, author = {Jixiao Zhang and Chunsheng Zuo}, year = {2025}, url = {https://arxiv.org/abs/2504.09696}, note = {Source identifier: 2504.09696} }