@misc{indiciaefcd98f40a075, title = {PR-OPD: Privileged Representation On-policy Self-Distillation for Agentic Reinforcement Learning}, author = {Muyang Li and Jie Yang and Zhengyu Fang and Junchao Zhu and Zhengkun Xiao and Ruining Deng and Zhe Jiang and Shigang Chen}, year = {2026}, url = {https://arxiv.org/abs/2609.36642}, note = {Source identifier: 2609.36642} }