@misc{indiciae5806cd623b1d, title = {Fairness Begins with State: Purifying Latent Preferences for Hierarchical Reinforcement Learning in Interactive Recommendation}, author = {Yun Lu and Xiaoyu Shi and Hong Xie and Xiangyu Zhao and Mingsheng Shang}, year = {2026}, url = {https://arxiv.org/abs/2603.03820}, note = {Source identifier: 2603.03820} }