@inproceedings{regnault-2026-gathering,
title = "Gathering valency frames for annotation and batch corrections",
author = "Regnault, Mathilde",
editor = {{\c{C}}{\"o}ltekin, {\c{C}}a{\u{g}}r{\i} and
Dobrovoljc, Kaja},
booktitle = "Proceedings of the Ninth Workshop on {U}niversal {D}ependencies ({UDW} 2026)",
month = may,
year = "2026",
address = "Palma de Mallorca, Spain",
publisher = "ELRA Language Resources Association (ELRA)",
url = "https://aclanthology.org/2026.udw-1.26/",
doi = "10.63317/3fcscsa4so7n",
pages = "289--294",
abstract = "Syntactic annotation is time- and resource-consuming, especially for historical and heterogeneous data. The Universal Dependencies (UD) framework provides a stable and cross-linguistically consistent annotation scheme, offering a crucial backbone for diachronic corpus studies. However, ensuring internal consistency within historical UD treebanks remains challenging due to syntactic variation and parser errors. We address this issue for Medieval and Classical French by integrating valency information into our corrections to support UD treebank maintenance. Valency frames were extracted from the Profiterole treebank (v. 2.7) and used to enrich OFrLex with structured valency information for Medieval French. Existing lexical resources such as Lefff are also exploited for Contemporary French. These valency frames are used to detect and correct inconsistencies in automatically annotated data through batch operations, thereby reinforcing UD guideline compliance and improving annotation coherence across diachronic stages. Preliminary experiments on Medieval French and exploratory annotation of Classical French data suggest that lexicon-informed error mining can reduce manual revision effort while strengthening the diachronic continuity enabled by the UD framework."
}<?xml version="1.0" encoding="UTF-8"?>
<modsCollection xmlns="http://www.loc.gov/mods/v3">
<mods ID="regnault-2026-gathering">
<titleInfo>
<title>Gathering valency frames for annotation and batch corrections</title>
</titleInfo>
<name type="personal">
<namePart type="given">Mathilde</namePart>
<namePart type="family">Regnault</namePart>
<role>
<roleTerm authority="marcrelator" type="text">author</roleTerm>
</role>
</name>
<originInfo>
<dateIssued>2026-05</dateIssued>
</originInfo>
<typeOfResource>text</typeOfResource>
<relatedItem type="host">
<titleInfo>
<title>Proceedings of the Ninth Workshop on Universal Dependencies (UDW 2026)</title>
</titleInfo>
<name type="personal">
<namePart type="given">Çağrı</namePart>
<namePart type="family">Çöltekin</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<name type="personal">
<namePart type="given">Kaja</namePart>
<namePart type="family">Dobrovoljc</namePart>
<role>
<roleTerm authority="marcrelator" type="text">editor</roleTerm>
</role>
</name>
<originInfo>
<publisher>ELRA Language Resources Association (ELRA)</publisher>
<place>
<placeTerm type="text">Palma de Mallorca, Spain</placeTerm>
</place>
</originInfo>
<genre authority="marcgt">conference publication</genre>
</relatedItem>
<abstract>Syntactic annotation is time- and resource-consuming, especially for historical and heterogeneous data. The Universal Dependencies (UD) framework provides a stable and cross-linguistically consistent annotation scheme, offering a crucial backbone for diachronic corpus studies. However, ensuring internal consistency within historical UD treebanks remains challenging due to syntactic variation and parser errors. We address this issue for Medieval and Classical French by integrating valency information into our corrections to support UD treebank maintenance. Valency frames were extracted from the Profiterole treebank (v. 2.7) and used to enrich OFrLex with structured valency information for Medieval French. Existing lexical resources such as Lefff are also exploited for Contemporary French. These valency frames are used to detect and correct inconsistencies in automatically annotated data through batch operations, thereby reinforcing UD guideline compliance and improving annotation coherence across diachronic stages. Preliminary experiments on Medieval French and exploratory annotation of Classical French data suggest that lexicon-informed error mining can reduce manual revision effort while strengthening the diachronic continuity enabled by the UD framework.</abstract>
<identifier type="citekey">regnault-2026-gathering</identifier>
<identifier type="doi">10.63317/3fcscsa4so7n</identifier>
<location>
<url>https://aclanthology.org/2026.udw-1.26/</url>
</location>
<part>
<date>2026-05</date>
<extent unit="page">
<start>289</start>
<end>294</end>
</extent>
</part>
</mods>
</modsCollection>
%0 Conference Proceedings
%T Gathering valency frames for annotation and batch corrections
%A Regnault, Mathilde
%Y Çöltekin, Çağrı
%Y Dobrovoljc, Kaja
%S Proceedings of the Ninth Workshop on Universal Dependencies (UDW 2026)
%D 2026
%8 May
%I ELRA Language Resources Association (ELRA)
%C Palma de Mallorca, Spain
%F regnault-2026-gathering
%X Syntactic annotation is time- and resource-consuming, especially for historical and heterogeneous data. The Universal Dependencies (UD) framework provides a stable and cross-linguistically consistent annotation scheme, offering a crucial backbone for diachronic corpus studies. However, ensuring internal consistency within historical UD treebanks remains challenging due to syntactic variation and parser errors. We address this issue for Medieval and Classical French by integrating valency information into our corrections to support UD treebank maintenance. Valency frames were extracted from the Profiterole treebank (v. 2.7) and used to enrich OFrLex with structured valency information for Medieval French. Existing lexical resources such as Lefff are also exploited for Contemporary French. These valency frames are used to detect and correct inconsistencies in automatically annotated data through batch operations, thereby reinforcing UD guideline compliance and improving annotation coherence across diachronic stages. Preliminary experiments on Medieval French and exploratory annotation of Classical French data suggest that lexicon-informed error mining can reduce manual revision effort while strengthening the diachronic continuity enabled by the UD framework.
%R 10.63317/3fcscsa4so7n
%U https://aclanthology.org/2026.udw-1.26/
%U https://doi.org/10.63317/3fcscsa4so7n
%P 289-294
Markdown (Informal)
[Gathering valency frames for annotation and batch corrections](https://aclanthology.org/2026.udw-1.26/) (Regnault, UDW 2026)
ACL