<?xml version="1.0" encoding="UTF-8"?>
<oai_dc:dc xmlns:oai_dc="http://www.openarchives.org/OAI/2.0/oai_dc/" xmlns:dc="http://purl.org/dc/elements/1.1/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/oai_dc/ http://www.openarchives.org/OAI/2.0/oai_dc.xsd">
  <dc:title>Automated REtrieval from TExt</dc:title>
  <dc:title>R package arete version 0.2</dc:title>
  <dc:description>A Python based pipeline for extraction of species occurrence data through the usage of large language models. Includes validation tools designed to handle model hallucinations for a scientific, rigorous use of LLM. Currently supports usage of GPT with more planned, including local and non-proprietary models. For more details on the methodology used please consult the references listed under each function, such as Kent, A. et al. (1995) &lt;doi:10.1002/asi.5090060209&gt;, van Rijsbergen, C.J. (1979, ISBN:978-0408709293, Levenshtein, V.I. (1966) &lt;https://nymity.ch/sybilhunting/pdf/Levenshtein1966a.pdf&gt; and Klaus Krippendorff (2011) &lt;https://repository.upenn.edu/handle/20.500.14332/2089&gt;.</dc:description>
  <dc:type>Software</dc:type>
  <dc:relation>Depends: R (&gt;= 4.3.0)</dc:relation>
  <dc:relation>Imports: terra, cld2, stringr, reticulate, pdftools, fedmatch,
kableExtra, dplyr, gecko, methods, ggplot2, jsonlite,
googledrive, irr, rmarkdown</dc:relation>
  <dc:relation>Suggests: knitr</dc:relation>
  <dc:creator>Vasco V. Branco &lt;vasco.branco@helsinki.fi&gt;</dc:creator>
  <dc:publisher>Comprehensive R Archive Network (CRAN)</dc:publisher>
  <dc:contributor>Vasco V. Branco [cre, aut] (ORCID:
    &lt;https://orcid.org/0000-0001-7797-3183&gt;),
  Vaughn Shirey [ctb] (ORCID: &lt;https://orcid.org/0000-0002-3589-9699&gt;),
  Thomas Merrien [ctb] (ORCID: &lt;https://orcid.org/0000-0002-0339-5656&gt;),
  Pedro Cardoso [aut] (ORCID: &lt;https://orcid.org/0000-0001-8119-9960&gt;)</dc:contributor>
  <dc:rights>GPL-3</dc:rights>
  <dc:date>2026-05-11</dc:date>
  <dc:format>application/tgz</dc:format>
  <dc:identifier>https://CRAN.R-project.org/package=arete</dc:identifier>
  <dc:identifier>doi:10.32614/CRAN.package.arete</dc:identifier>
</oai_dc:dc>
