@misc{indiciaefc9b59e6276d, title = {No Train No Gain: Revisiting Efficient Training Algorithms For Transformer-based Language Models}, author = {Jean Kaddour and Oscar Key and Piotr Nawrot and Pasquale Minervini and Matt J. Kusner}, year = {2023}, url = {https://arxiv.org/abs/2307.06440}, note = {Source identifier: 2307.06440} }