@misc{indiciae71f5a0eebc0c, title = {Train Flat, Then Compress: Sharpness-Aware Minimization Learns More Compressible Models}, author = {Clara Na and Sanket Vaibhav Mehta and Emma Strubell}, year = {2022}, doi = {10.18653/v1/2022.findings-emnlp.361}, url = {https://arxiv.org/abs/2205.12694}, note = {Source identifier: 2205.12694} }