<?xml version='1.0' encoding='UTF-8'?><?xml-stylesheet href='static/style.xsl' type='text/xsl'?><OAI-PMH xmlns="http://www.openarchives.org/OAI/2.0/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/ http://www.openarchives.org/OAI/2.0/OAI-PMH.xsd"><responseDate>2026-08-09T13:43:40Z</responseDate><request verb="GetRecord" identifier="oai:www.clarin.si:11356/1756" metadataPrefix="oai_dc">http://www.clarin.si/repository/oai/request</request><GetRecord><record><header><identifier>oai:www.clarin.si:11356/1756</identifier><datestamp>2024-11-06T17:20:40Z</datestamp><setSpec>hdl_11356_1023</setSpec><setSpec>hdl_11356_1024</setSpec></header><metadata><oai_dc:dc xmlns:oai_dc="http://www.openarchives.org/OAI/2.0/oai_dc/" xmlns:doc="http://www.lyncode.com/xoai" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xmlns:dc="http://purl.org/dc/elements/1.1/" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/oai_dc/ http://www.openarchives.org/OAI/2.0/oai_dc.xsd">
<dc:title>Slovene translation of the SQuAD2.0 dataset</dc:title>
<dc:creator>Borovič, Mladen</dc:creator>
<dc:creator>Žagar, Kristjan</dc:creator>
<dc:creator>Ferme, Marko</dc:creator>
<dc:creator>Majninger, Sandi</dc:creator>
<dc:creator>Ojsteršek, Milan</dc:creator>
<dc:creator>Šmajdek, Uroš</dc:creator>
<dc:creator>Zirkelbach, Maj</dc:creator>
<dc:creator>Zupanič, Matjaž</dc:creator>
<dc:creator>Jazbinšek, Meta</dc:creator>
<dc:creator>Žitnik, Slavko</dc:creator>
<dc:creator>Robnik-Šikonja, Marko</dc:creator>
<dc:subject>Q&amp;A</dc:subject>
<dc:subject>SQuAD</dc:subject>
<dc:subject>natural language processing</dc:subject>
<dc:description>Stanford Question Answering Dataset (SQuAD) is a reading comprehension dataset, consisting of questions posed by crowdworkers on a set of Wikipedia articles, where the answer to every question is a segment of text, or span, from the corresponding reading passage, or the question might be unanswerable. SQuAD2.0 combines the 100,000 questions in SQuAD1.1 with over 50,000 unanswerable questions written adversarially by crowdworkers to look similar to answerable ones. To do well on SQuAD2.0, systems must not only answer questions when possible, but also determine when no answer is supported by the paragraph and abstain from answering. The English version of SQuAD2.0 was machine translated to Slovene, then the translation was manually reviewed and corrected where needed. The data is provided in JSON format and consists of a training set and a validation set.</dc:description>
<dc:date>2022-09-22</dc:date>
<dc:type>corpus</dc:type>
<dc:identifier>http://hdl.handle.net/11356/1756</dc:identifier>
<dc:language>slv</dc:language>
<dc:relation>https://rajpurkar.github.io/SQuAD-explorer/</dc:relation>
<dc:rights>Creative Commons - Attribution 4.0 International (CC BY 4.0)</dc:rights>
<dc:rights>https://creativecommons.org/licenses/by/4.0/</dc:rights>
<dc:rights>PUB</dc:rights>
<dc:format>text/plain; charset=utf-8</dc:format>
<dc:format>application/octet-stream</dc:format>
<dc:format>application/octet-stream</dc:format>
<dc:format>downloadable_files_count: 2</dc:format>
<dc:publisher>Faculty of Electrical Engineering and Computer Science, University of Maribor</dc:publisher>
<dc:publisher>Faculty of Computer and Information Science, University of Ljubljana</dc:publisher>
<dc:publisher>Faculty of Arts, University of Ljubljana</dc:publisher>
<dc:source>https://rsdo.slovenscina.eu/</dc:source>
</oai_dc:dc>
</metadata></record></GetRecord></OAI-PMH>