@misc{indiciae1479b736b389, title = {Turning Trash into Treasure: Accelerating Inference of Large Language Models with Token Recycling}, author = {Xianzhen Luo and Yixuan Wang and Qingfu Zhu and Zhiming Zhang and Xuanyu Zhang and Qing Yang and Dongliang Xu}, year = {2025}, url = {https://arxiv.org/abs/2408.08696}, note = {Source identifier: 2408.08696} }