@misc{indiciae210fc342095a, title = {DepthKV: Layer-Dependent KV Cache Pruning for Long-Context LLM Inference}, author = {Zahra Dehghanighobadi and Asja Fischer}, year = {2026}, url = {https://arxiv.org/abs/2604.24647}, note = {Source identifier: 2604.24647} }