Views
No views yet
| Component | Description |
|---|---|
| Encoder | XLM-RoBERTa base (CINO v2 variant) |
| Decoder | Hybrid transformer with: |
| • NormalDecoderLayer: Randomly initialized standard layers | |
| • CustomDecoderLayer: Weight-shared layers with dual FFN structure | |
| Parameters | 492M total parameters |
@article{su2025multilingualencoderknowsrealize,
author = {Zeli Su and Ziyin Zhang and Guixian Xu and Jianing Liu and Xu Han and Ting Zhang and Yushuang Dong},
title = {Multilingual Encoder Knows more than You Realize: Shared Weights Pretraining
for Extremely Low-Resource Languages},
journal = {CoRR},
volume = {abs/2502.10852},
year = {2025},
url = {https://doi.org/10.48550/arXiv.2502.10852},
doi = {10.48550/ARXIV.2502.10852},
eprinttype = {arXiv},
eprint = {2502.10852}
}