<?xml version="1.0" encoding="UTF-8"?><?xml-stylesheet type="text/xsl" href="static/style.xsl"?><OAI-PMH xmlns="http://www.openarchives.org/OAI/2.0/" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xsi:schemaLocation="http://www.openarchives.org/OAI/2.0/ http://www.openarchives.org/OAI/2.0/OAI-PMH.xsd"><responseDate>2026-09-16T03:49:51Z</responseDate><request verb="GetRecord" identifier="oai:gredos.usal.es:10366/56126" metadataPrefix="marc">https://gredos.usal.es/oai/request</request><GetRecord><record><header><identifier>oai:gredos.usal.es:10366/56126</identifier><datestamp>2025-06-05T12:40:28Z</datestamp><setSpec>com_10366_4558</setSpec><setSpec>com_10366_4512</setSpec><setSpec>com_10366_3823</setSpec><setSpec>col_10366_4559</setSpec></header><metadata><record xmlns="http://www.loc.gov/MARC21/slim" xmlns:doc="http://www.lyncode.com/xoai" xmlns:xsi="http://www.w3.org/2001/XMLSchema-instance" xmlns:dcterms="http://purl.org/dc/terms/" xsi:schemaLocation="http://www.loc.gov/MARC21/slim http://www.loc.gov/standards/marcxml/schema/MARC21slim.xsd">
<leader>00925njm 22002777a 4500</leader>
<datafield tag="042" ind1=" " ind2=" ">
<subfield code="a">dc</subfield>
</datafield>
<datafield tag="720" ind1=" " ind2=" ">
<subfield code="a">García de Figuerola Paniagua, Luis Carlos</subfield>
<subfield code="e">author</subfield>
</datafield>
<datafield tag="720" ind1=" " ind2=" ">
<subfield code="a">Gómez Díaz, Raquel</subfield>
<subfield code="e">author</subfield>
</datafield>
<datafield tag="720" ind1=" " ind2=" ">
<subfield code="a">López de San Román, Eva</subfield>
<subfield code="e">author</subfield>
</datafield>
<datafield tag="260" ind1=" " ind2=" ">
<subfield code="c">2000</subfield>
</datafield>
<datafield tag="520" ind1=" " ind2=" ">
<subfield code="a">At some stage, most of the models and techniques implemented in IR use frequency counts of the terms appearing in documents and in queries. However, many words, since they are derived from the same stem, have very close semantic contents. This makes a grouping of such variants under a single term advisable. Otherwise, dispersal occurs in the calculation of frequency of these terms, and it also becomes difficult to compare queries and documents. On the other hand, there are notable differences between different languages in the way of forming derivatives and inflected forms, so that the application of specific techniques can produce unequal results according to the language of the documents and queries. A description is given of the tests carried out for documents in Spanish, which involved some stemming techniques widely used in English, as well as the application of n-grams, and the results are compared.</subfield>
</datafield>
<datafield tag="024" ind2=" " ind1="8">
<subfield code="a">Figuerola García, L.C., Gómez Díaz, R. y López de San Román, E. (2000). Stemming and n-grams in Spanish: an evaluation of their impact on information retrieval. "Journal of Information Science",  26 (6), 461-467.</subfield>
</datafield>
<datafield tag="024" ind2=" " ind1="8">
<subfield code="a">http://hdl.handle.net/10366/56126</subfield>
</datafield>
<datafield ind1=" " ind2=" " tag="653">
<subfield code="a">Recuperación de la información</subfield>
</datafield>
<datafield ind1=" " ind2=" " tag="653">
<subfield code="a">S-stemmer</subfield>
</datafield>
<datafield ind1=" " ind2=" " tag="653">
<subfield code="a">N-gramas</subfield>
</datafield>
<datafield ind1=" " ind2=" " tag="653">
<subfield code="a">Information retrieval</subfield>
</datafield>
<datafield ind1=" " ind2=" " tag="653">
<subfield code="a">Stemming</subfield>
</datafield>
<datafield ind1=" " ind2=" " tag="653">
<subfield code="a">N-grams</subfield>
</datafield>
<datafield tag="245" ind1="0" ind2="0">
<subfield code="a">Stemming and n-grams in Spanish: an evaluation of their impact on information retrieval</subfield>
</datafield>
</record></metadata></record></GetRecord></OAI-PMH>