@misc{indiciae129eae7e4cee, title = {How to Compress KV Cache in RL Post-Training? Shadow Mask Distillation for Memory-Efficient Alignment}, author = {Rui Zhu and Weiheng Bai and Qiushi Wu and Yang Ren and Haixu Tang and Yuchu Liu}, year = {2026}, url = {https://arxiv.org/abs/2605.06850}, note = {Source identifier: 2605.06850} }