@misc{indiciae3d1af64da208, title = {RL's Razor: Why Online Reinforcement Learning Forgets Less}, author = {Idan Shenfeld and Jyothish Pari and Pulkit Agrawal}, year = {2025}, url = {https://arxiv.org/abs/2509.04259}, note = {Source identifier: 2509.04259} }