@misc{indiciae78195e0a3ef0, title = {H1B-KV: Hybrid One-Bit Caches for Memory-Efficient Large Language Model Inference}, author = {Harshil Vejendla}, year = {2025}, url = {https://arxiv.org/abs/2510.05529}, note = {Source identifier: 2510.05529} }