@inproceedings{WartenaSommer2012, author = {Christian Wartena and Maike Sommer}, title = {Automatic classification of scientific records using the German Subject Heading Authority File (SWD)}, series = {Proceedings of the 2nd International Workshop on Semantic Digital Archives (SDA 2012)}, doi = {10.25968/opus-328}, url = {http://nbn-resolving.de/urn:nbn:de:bsz:960-opus-4008}, pages = {37 -- 48}, year = {2012}, abstract = {The following paper deals with an automatic text classification method which does not require training documents. For this method the German Subject Heading Authority File (SWD), provided by the linked data service of the German National Library is used. Recently the SWD was enriched with notations of the Dewey Decimal Classification (DDC). In consequence it became possible to utilize the subject headings as textual representations for the notations of the DDC. Basically, we we derive the classification of a text from the classification of the words in the text given by the thesaurus. The method was tested by classifying 3826 OAI-Records from 7 different repositories. Mean reciprocal rank and recall were chosen as evaluation measure. Direct comparison to a machine learning method has shown that this method is definitely competitive. Thus we can conclude that the enriched version of the SWD provides high quality information with a broad coverage for classification of German scientific articles.}, language = {en} }