@misc{indiciae3c7c645e9aef, title = {The Missing Piece in Pre-trained Model Evaluation: Reward-Guided Decoding Unlocks Task-Oriented Behavior Without Parameter Updates}, author = {Shaobo Wang and Guo Chen and Ziyue Wang and Zhengyang Tang and Qingyang Liu and Xingzhang Ren and Dayiheng Liu and Linfeng Zhang}, year = {2026}, url = {https://arxiv.org/abs/2605.28020}, note = {Source identifier: 2605.28020} }