@misc{indiciaeb93496a68b05, title = {Efficient LLM Inference with Activation Checkpointing and Hybrid Caching}, author = {Sanghyeon Lee and Hongbeen Kim and Soojin Hwang and Guseul Heo and Minwoo Noh and Jaehyuk Huh}, year = {2025}, doi = {10.1109/iccd65941.2025.00045}, url = {https://arxiv.org/abs/2501.01792}, note = {Source identifier: 2501.01792} }