@misc{indiciae94514673f5cf, title = {Shfl-BW: Accelerating Deep Neural Network Inference with Tensor-Core Aware Weight Pruning}, author = {Guyue Huang and Haoran Li and Minghai Qin and Fei Sun and Yufei Ding and Yuan Xie}, year = {2022}, url = {https://arxiv.org/abs/2203.05016}, note = {Source identifier: 2203.05016} }