@misc{indiciaed966010cdf72, title = {L1: Controlling How Long A Reasoning Model Thinks With Reinforcement Learning}, author = {Pranjal Aggarwal and Sean Welleck}, year = {2025}, url = {https://arxiv.org/abs/2503.04697}, note = {Source identifier: 2503.04697} }