@misc{indiciae72e4f2a491d9, title = {Representational alignment yields generalizable safety in language models}, author = {Lingyu Li and Yan Teng and Yingchun Wang and Xia Hu}, year = {2026}, url = {https://arxiv.org/abs/2609.04022}, note = {Source identifier: 2609.04022} }