@misc{indiciae5b545ecf963e, title = {CachePrune: Privacy-Aware and Fine-Grained KV Cache Sharing for Efficient LLM Inference}, author = {Guanlong Wu and Zhaohan li and Yao Zhang and Zheng Zhang and Jianyu Niu and Ye Wu and Yinqian Zhang}, year = {2026}, url = {https://arxiv.org/abs/2605.23640}, note = {Source identifier: 2605.23640} }