@misc{indiciae5183367daa8a, title = {VL-Cache: Sparsity and Modality-Aware KV Cache Compression for Vision-Language Model Inference Acceleration}, author = {Dezhan Tu and Danylo Vashchilenko and Yuzhe Lu and Panpan Xu}, year = {2024}, url = {https://arxiv.org/abs/2410.23317}, note = {Source identifier: 2410.23317} }