@misc{indiciae6696d1c699fb, title = {Accelerating Transformer Inference and Training with 2:4 Activation Sparsity}, author = {Daniel Haziza and Timothy Chou and Dhruv Choudhary and Luca Wehrstedt and Francisco Massa and Jiecao Yu and Geonhwa Jeong and Supriya Rao and Patrick Labatut and Jesse Cai}, year = {2025}, url = {https://arxiv.org/abs/2503.16672}, note = {Source identifier: 2503.16672} }