@misc{indiciae762d68bc4358, title = {Training neural networks faster with minimal tuning using pre-computed lists of hyperparameters for NAdamW}, author = {Sourabh Medapati and Priya Kasimbeg and Shankar Krishnan and Naman Agarwal and George Dahl}, year = {2025}, url = {https://arxiv.org/abs/2503.03986}, note = {Source identifier: 2503.03986} }