@misc{indiciae5ede6cc38dfb, title = {Reinforcement Learning Towards Broadly and Persistently Beneficial Models}, author = {Akshay V. Jagadeesh and Rahul K. Arora and Khaled Saab and Ali Malik and Mikhail Trofimov and Foivos Tsimpourlas and Johannes Heidecke and Karan Singhal}, year = {2026}, url = {https://arxiv.org/abs/2606.24014}, note = {Source identifier: 2606.24014} }