@misc{indiciae4b355348ea7e, title = {Reducing the Probability of Undesirable Outputs in Language Models Using Probabilistic Inference}, author = {Stephen Zhao and Aidan Li and Rob Brekelmans and Roger Grosse}, year = {2025}, url = {https://arxiv.org/abs/2510.21184}, note = {Source identifier: 2510.21184} }