@misc{indiciaefcfc57a66403, title = {Ps and Qs: Quantization-aware pruning for efficient low latency neural network inference}, author = {Benjamin Hawks and Javier Duarte and Nicholas J. Fraser and Alessandro Pappalardo and Nhan Tran and Yaman Umuroglu}, year = {2021}, doi = {10.3389/frai.2021.676564}, url = {https://arxiv.org/abs/2102.11289}, note = {Source identifier: 2102.11289} }