@misc{indiciae42f86c4e902d, title = {Faithful-Patchscopes: Understanding and Mitigating Model Bias in Hidden Representations Explanation of Large Language Models}, author = {Xilin Gong and Shu Yang and Zehua Cao and Lynne Billard and Di Wang}, year = {2026}, url = {https://arxiv.org/abs/2602.00300}, note = {Source identifier: 2602.00300} }