@misc{indiciae57a6b749585d, title = {Lookahead: An Inference Acceleration Framework for Large Language Model with Lossless Generation Accuracy}, author = {Yao Zhao and Zhitian Xie and Chen Liang and Chenyi Zhuang and Jinjie Gu}, year = {2024}, url = {https://arxiv.org/abs/2312.12728}, note = {Source identifier: 2312.12728} }