@misc{indiciaecb120d798a33, title = {Lossless Acceleration of Large Language Model via Adaptive N-gram Parallel Decoding}, author = {Jie Ou and Yueming Chen and Wenhong Tian}, year = {2024}, url = {https://arxiv.org/abs/2404.08698}, note = {Source identifier: 2404.08698} }