@misc{indiciae7e70c21bb715, title = {When Data Is Scarce: Scaling Sparse Language Models with Repeated Training}, author = {Boqian Wu and Qiao Xiao and Patrik Okanovic and Tomasz Sternal and Maurice van Keulen and Mykola Pechenizkiy and Elena Mocanu and Torsten Hoefler and Decebal Constantin Mocanu}, year = {2026}, url = {https://arxiv.org/abs/2606.01155}, note = {Source identifier: 2606.01155} }