@inproceedings{campos-etal-2023-oberta,
title = "o{BERT}a: Improving Sparse Transfer Learning via improved initialization, distillation, and pruning regimes",
author = "Campos, Daniel and
Marques, Alexandre and
Kurtz, Mark and
Xiang Zhai, Cheng",
editor = "Sadat Moosavi, Nafise and
Gurevych, Iryna and
Hou, Yufang and
Kim, Gyuwan and
Kim, Young Jin and
Schuster, Tal and
Agrawal, Ameeta",
booktitle = "Proceedings of The Fourth Workshop on Simple and Efficient Natural Language Processing (SustaiNLP)",
month = jul,
year = "2023",
address = "Toronto, Canada (Hybrid)",
publisher = "Association for Computational Linguistics",
url = "https://aclanthology.org/2023.sustainlp-1.3",
doi = "10.18653/v1/2023.sustainlp-1.3",
pages = "39--58",
}
<?xml version="1.0" encoding="UTF-8"?>
<modsCollection xmlns="http://www.loc.gov/mods/v3">
<mods ID="campos-etal-2023-oberta">
<titleInfo>
<title>oBERTa: Improving Sparse Transfer Learning via improved initialization, distillation, and pruning regimes</title>
</titleInfo>
<name type="personal">
<namePart type="given">Daniel</namePart>
<namePart type="family">Campos</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Alexandre</namePart>
<namePart type="family">Marques</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Mark</namePart>
<namePart type="family">Kurtz</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Cheng</namePart>
<namePart type="family">Xiang Zhai</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<originInfo>
<dateIssued>2023-07</dateIssued>
</originInfo>
<typeOfResource>text</typeOfResource>
<relatedItem type="host">
<titleInfo>
<title>Proceedings of The Fourth Workshop on Simple and Efficient Natural Language Processing (SustaiNLP)</title>
</titleInfo>
<name type="personal">
<namePart type="given">Nafise</namePart>
<namePart type="family">Sadat Moosavi</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Iryna</namePart>
<namePart type="family">Gurevych</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Yufang</namePart>
<namePart type="family">Hou</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Gyuwan</namePart>
<namePart type="family">Kim</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Young</namePart>
<namePart type="given">Jin</namePart>
<namePart type="family">Kim</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Tal</namePart>
<namePart type="family">Schuster</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Ameeta</namePart>
<namePart type="family">Agrawal</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<originInfo>
<publisher>Association for Computational Linguistics</publisher>
<place>
<placeTerm type="text">Toronto, Canada (Hybrid)</placeTerm>
</place>
</originInfo>
<genre authority="marcgt">conference publication</genre>
</relatedItem>
<identifier type="citekey">campos-etal-2023-oberta</identifier>
<identifier type="doi">10.18653/v1/2023.sustainlp-1.3</identifier>
<location>
<url>https://aclanthology.org/2023.sustainlp-1.3</url>
</location>
<part>
<date>2023-07</date>
<extent unit="page">
<start>39</start>
<end>58</end>
</extent>
</part>
</mods>
</modsCollection>
%0 Conference Proceedings
%T oBERTa: Improving Sparse Transfer Learning via improved initialization, distillation, and pruning regimes
%A Campos, Daniel
%A Marques, Alexandre
%A Kurtz, Mark
%A Xiang Zhai, Cheng
%Y Sadat Moosavi, Nafise
%Y Gurevych, Iryna
%Y Hou, Yufang
%Y Kim, Gyuwan
%Y Kim, Young Jin
%Y Schuster, Tal
%Y Agrawal, Ameeta
%S Proceedings of The Fourth Workshop on Simple and Efficient Natural Language Processing (SustaiNLP)
%D 2023
%8 July
%I Association for Computational Linguistics
%C Toronto, Canada (Hybrid)
%F campos-etal-2023-oberta
%R 10.18653/v1/2023.sustainlp-1.3
%U https://aclanthology.org/2023.sustainlp-1.3
%U https://doi.org/10.18653/v1/2023.sustainlp-1.3
%P 39-58
Markdown (Informal)
[oBERTa: Improving Sparse Transfer Learning via improved initialization, distillation, and pruning regimes](https://aclanthology.org/2023.sustainlp-1.3) (Campos et al., sustainlp 2023)
ACL