@misc{indiciae1d2b8e2e346b, title = {Do pretrained Transformers Learn In-Context by Gradient Descent?}, author = {Lingfeng Shen and Aayush Mishra and Daniel Khashabi}, year = {2024}, url = {https://arxiv.org/abs/2310.08540}, note = {Source identifier: 2310.08540} }