@misc{indiciae20f3597c1487, title = {GPTQT: Quantize Large Language Models Twice to Push the Efficiency}, author = {Yipin Guo and Yilin Lang and Qinyuan Ren}, year = {2024}, url = {https://arxiv.org/abs/2407.02891}, note = {Source identifier: 2407.02891} }