@misc{indiciae30186209b2a8, title = {Dynamic Rewarding with Prompt Optimization Enables Tuning-free Self-Alignment of Language Models}, author = {Somanshu Singla and Zhen Wang and Tianyang Liu and Abdullah Ashfaq and Zhiting Hu and Eric P. Xing}, year = {2024}, url = {https://arxiv.org/abs/2411.08733}, note = {Source identifier: 2411.08733} }