@misc{indiciaebbced444a330, title = {ABQ-LLM: Arbitrary-Bit Quantized Inference Acceleration for Large Language Models}, author = {Chao Zeng and Songwei Liu and Yusheng Xie and Hong Liu and Xiaojian Wang and Miao Wei and Shu Yang and Fangmin Chen and Xing Mei}, year = {2025}, url = {https://arxiv.org/abs/2408.08554}, note = {Source identifier: 2408.08554} }