@misc{indiciae7d0656a75279, title = {h1: Bootstrapping LLMs to Reason over Longer Horizons via Reinforcement Learning}, author = {Sumeet Ramesh Motwani and Alesia Ivanova and Ziyang Cai and Philip Torr and Riashat Islam and Shital Shah and Christian Schroeder de Witt and Charles London}, year = {2025}, url = {https://arxiv.org/abs/2510.07312}, note = {Source identifier: 2510.07312} }