@misc{indiciaefe2d990dd2b2, title = {Fast-dLLM: Training-free Acceleration of Diffusion LLM by Enabling KV Cache and Parallel Decoding}, author = {Chengyue Wu and Hao Zhang and Shuchen Xue and Zhijian Liu and Shizhe Diao and Ligeng Zhu and Ping Luo and Song Han and Enze Xie}, year = {2025}, url = {https://arxiv.org/abs/2505.22618}, note = {Source identifier: 2505.22618} }