@misc{indiciae56e1bb57326a, title = {On the Optimal Reasoning Length for RL-Trained Language Models}, author = {Daisuke Nohara and Taishi Nakamura and Rio Yokota}, year = {2026}, url = {https://arxiv.org/abs/2602.09591}, note = {Source identifier: 2602.09591} }