@misc{indiciae3acaa3ff26ec, title = {Mitigating Overthinking in Large Reasoning Models via Difficulty-aware Reinforcement Learning}, author = {Qian Wan and Ziao Xu and Luona Wei and Xiaoxuan Shen and Jianwen Sun}, year = {2026}, url = {https://arxiv.org/abs/2601.21418}, note = {Source identifier: 2601.21418} }