@misc{indiciae5d41f4364b81, title = {Metacognition as Reward: Reinforcing LLM Reasoning via Knowledge and Regulation Signals}, author = {Sirui Chen and Lei Xu and Yuying Zhao and Yutian Chen and Yu Wang and Beier Zhu and Hanwang Zhang and Shengjie Zhao and Chaochao Lu}, year = {2026}, url = {https://arxiv.org/abs/2605.23384}, note = {Source identifier: 2605.23384} }