@misc{indiciae7dd40e56c586, title = {The Past, Present and Better Future of Feedback Learning in Large Language Models for Subjective Human Preferences and Values}, author = {Hannah Rose Kirk and Andrew M. Bean and Bertie Vidgen and Paul Röttger and Scott A. Hale}, year = {2023}, url = {https://arxiv.org/abs/2310.07629}, note = {Source identifier: 2310.07629} }