@misc{indiciae4e7bed4225d0, title = {Training Dynamics Underlying Language Model Scaling Laws: Loss Deceleration and Zero-Sum Learning}, author = {Andrei Mircea and Supriyo Chakraborty and Nima Chitsazan and Milind Naphade and Sambit Sahu and Irina Rish and Ekaterina Lobacheva}, year = {2025}, url = {https://arxiv.org/abs/2506.05447}, note = {Source identifier: 2506.05447} }