{"id":"https://openalex.org/W2290619535","doi":"https://doi.org/10.1007/978-3-319-25252-0_46","title":"Harvesting Comparable Corpora and Mining Them for Equivalent Bilingual Sentences Using Statistical Classification and Analogy-Based Heuristics","display_name":"Harvesting Comparable Corpora and Mining Them for Equivalent Bilingual Sentences Using Statistical Classification and Analogy-Based Heuristics","publication_year":2015,"publication_date":"2015-01-01","ids":{"openalex":"https://openalex.org/W2290619535","doi":"https://doi.org/10.1007/978-3-319-25252-0_46","mag":"2290619535"},"language":"en","primary_location":{"is_oa":false,"landing_page_url":"https://doi.org/10.1007/978-3-319-25252-0_46","pdf_url":null,"source":{"id":"https://openalex.org/S106296714","display_name":"Lecture notes in computer science","issn_l":"0302-9743","issn":["0302-9743","1611-3349"],"is_oa":false,"is_in_doaj":false,"is_indexed_in_scopus":true,"is_core":true,"host_organization":"https://openalex.org/P4310319900","host_organization_name":"Springer Science+Business Media","host_organization_lineage":["https://openalex.org/P4310319965","https://openalex.org/P4310319900"],"host_organization_lineage_names":["Springer Nature","Springer Science+Business Media"],"type":"book series"},"license":null,"license_id":null,"version":null,"is_accepted":false,"is_published":false},"type":"book-chapter","type_crossref":"book-chapter","indexed_in":["arxiv","crossref","datacite"],"open_access":{"is_oa":true,"oa_status":"green","oa_url":"http://arxiv.org/pdf/1511.06285","any_repository_has_fulltext":true},"authorships":[{"author_position":"first","author":{"id":"https://openalex.org/A5030149703","display_name":"Krzysztof Wo\u0142k","orcid":"https://orcid.org/0000-0001-5030-334X"},"institutions":[{"id":"https://openalex.org/I3017851245","display_name":"Polish-Japanese Academy of Information Technology","ror":"https://ror.org/01v542j61","country_code":"PL","type":"education","lineage":["https://openalex.org/I3017851245"]}],"countries":["PL"],"is_corresponding":true,"raw_author_name":"Krzysztof Wo\u0142k","raw_affiliation_strings":["Department of Multimedia, Polish - Japanese Academy of Information Technology, Warsaw, Poland"],"affiliations":[{"raw_affiliation_string":"Department of Multimedia, Polish - Japanese Academy of Information Technology, Warsaw, Poland","institution_ids":["https://openalex.org/I3017851245"]}]},{"author_position":"middle","author":{"id":"https://openalex.org/A5088186005","display_name":"Emilia Rejmund","orcid":null},"institutions":[{"id":"https://openalex.org/I3017851245","display_name":"Polish-Japanese Academy of Information Technology","ror":"https://ror.org/01v542j61","country_code":"PL","type":"education","lineage":["https://openalex.org/I3017851245"]}],"countries":["PL"],"is_corresponding":false,"raw_author_name":"Emilia Rejmund","raw_affiliation_strings":["Department of Multimedia, Polish - Japanese Academy of Information Technology, Warsaw, Poland"],"affiliations":[{"raw_affiliation_string":"Department of Multimedia, Polish - Japanese Academy of Information Technology, Warsaw, Poland","institution_ids":["https://openalex.org/I3017851245"]}]},{"author_position":"last","author":{"id":"https://openalex.org/A5031351366","display_name":"Krzysztof Marasek","orcid":"https://orcid.org/0000-0003-1344-3524"},"institutions":[{"id":"https://openalex.org/I3017851245","display_name":"Polish-Japanese Academy of Information Technology","ror":"https://ror.org/01v542j61","country_code":"PL","type":"education","lineage":["https://openalex.org/I3017851245"]}],"countries":["PL"],"is_corresponding":false,"raw_author_name":"Krzysztof Marasek","raw_affiliation_strings":["Department of Multimedia, Polish - Japanese Academy of Information Technology, Warsaw, Poland"],"affiliations":[{"raw_affiliation_string":"Department of Multimedia, Polish - Japanese Academy of Information Technology, Warsaw, Poland","institution_ids":["https://openalex.org/I3017851245"]}]}],"institution_assertions":[],"countries_distinct_count":1,"institutions_distinct_count":1,"corresponding_author_ids":["https://openalex.org/A5030149703"],"corresponding_institution_ids":["https://openalex.org/I3017851245"],"apc_list":{"value":5000,"currency":"EUR","value_usd":5392},"apc_paid":null,"fwci":0.304,"has_fulltext":true,"fulltext_origin":"pdf","cited_by_count":2,"citation_normalized_percentile":{"value":0.5204,"is_in_top_1_percent":false,"is_in_top_10_percent":false},"cited_by_percentile_year":{"min":73,"max":76},"biblio":{"volume":null,"issue":null,"first_page":"433","last_page":"441"},"is_retracted":false,"is_paratext":false,"primary_topic":{"id":"https://openalex.org/T10181","display_name":"Natural Language Processing Techniques","score":0.9999,"subfield":{"id":"https://openalex.org/subfields/1702","display_name":"Artificial Intelligence"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}},"topics":[{"id":"https://openalex.org/T10181","display_name":"Natural Language Processing Techniques","score":0.9999,"subfield":{"id":"https://openalex.org/subfields/1702","display_name":"Artificial Intelligence"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}},{"id":"https://openalex.org/T10028","display_name":"Topic Modeling","score":0.9953,"subfield":{"id":"https://openalex.org/subfields/1702","display_name":"Artificial Intelligence"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}},{"id":"https://openalex.org/T12380","display_name":"Authorship Attribution and Profiling","score":0.9898,"subfield":{"id":"https://openalex.org/subfields/1702","display_name":"Artificial Intelligence"},"field":{"id":"https://openalex.org/fields/17","display_name":"Computer Science"},"domain":{"id":"https://openalex.org/domains/3","display_name":"Physical Sciences"}}],"keywords":[{"id":"https://openalex.org/keywords/heuristics","display_name":"Heuristics","score":0.7048778},{"id":"https://openalex.org/keywords/parallel-corpora","display_name":"Parallel corpora","score":0.5049558},{"id":"https://openalex.org/keywords/crawling","display_name":"Crawling","score":0.45827436}],"concepts":[{"id":"https://openalex.org/C41008148","wikidata":"https://www.wikidata.org/wiki/Q21198","display_name":"Computer science","level":0,"score":0.81470394},{"id":"https://openalex.org/C204321447","wikidata":"https://www.wikidata.org/wiki/Q30642","display_name":"Natural language processing","level":1,"score":0.71614337},{"id":"https://openalex.org/C127705205","wikidata":"https://www.wikidata.org/wiki/Q5748245","display_name":"Heuristics","level":2,"score":0.7048778},{"id":"https://openalex.org/C154945302","wikidata":"https://www.wikidata.org/wiki/Q11660","display_name":"Artificial intelligence","level":1,"score":0.6273044},{"id":"https://openalex.org/C203005215","wikidata":"https://www.wikidata.org/wiki/Q79798","display_name":"Machine translation","level":2,"score":0.6186088},{"id":"https://openalex.org/C2777855551","wikidata":"https://www.wikidata.org/wiki/Q12310021","display_name":"Subject (documents)","level":2,"score":0.53190255},{"id":"https://openalex.org/C2780451532","wikidata":"https://www.wikidata.org/wiki/Q759676","display_name":"Task (project management)","level":2,"score":0.51549155},{"id":"https://openalex.org/C2985367798","wikidata":"https://www.wikidata.org/wiki/Q1346592","display_name":"Parallel corpora","level":3,"score":0.5049558},{"id":"https://openalex.org/C521332185","wikidata":"https://www.wikidata.org/wiki/Q185816","display_name":"Analogy","level":2,"score":0.49487585},{"id":"https://openalex.org/C100368936","wikidata":"https://www.wikidata.org/wiki/Q1411725","display_name":"Crawling","level":2,"score":0.45827436},{"id":"https://openalex.org/C23123220","wikidata":"https://www.wikidata.org/wiki/Q816826","display_name":"Information retrieval","level":1,"score":0.42140883},{"id":"https://openalex.org/C206345919","wikidata":"https://www.wikidata.org/wiki/Q20380951","display_name":"Resource (disambiguation)","level":2,"score":0.41505468},{"id":"https://openalex.org/C136764020","wikidata":"https://www.wikidata.org/wiki/Q466","display_name":"World Wide Web","level":1,"score":0.18148014},{"id":"https://openalex.org/C41895202","wikidata":"https://www.wikidata.org/wiki/Q8162","display_name":"Linguistics","level":1,"score":0.15495479},{"id":"https://openalex.org/C127413603","wikidata":"https://www.wikidata.org/wiki/Q11023","display_name":"Engineering","level":0,"score":0.060375422},{"id":"https://openalex.org/C138885662","wikidata":"https://www.wikidata.org/wiki/Q5891","display_name":"Philosophy","level":0,"score":0.0},{"id":"https://openalex.org/C105702510","wikidata":"https://www.wikidata.org/wiki/Q514","display_name":"Anatomy","level":1,"score":0.0},{"id":"https://openalex.org/C111919701","wikidata":"https://www.wikidata.org/wiki/Q9135","display_name":"Operating system","level":1,"score":0.0},{"id":"https://openalex.org/C31258907","wikidata":"https://www.wikidata.org/wiki/Q1301371","display_name":"Computer network","level":1,"score":0.0},{"id":"https://openalex.org/C201995342","wikidata":"https://www.wikidata.org/wiki/Q682496","display_name":"Systems engineering","level":1,"score":0.0},{"id":"https://openalex.org/C71924100","wikidata":"https://www.wikidata.org/wiki/Q11190","display_name":"Medicine","level":0,"score":0.0}],"mesh":[],"locations_count":4,"locations":[{"is_oa":false,"landing_page_url":"https://doi.org/10.1007/978-3-319-25252-0_46","pdf_url":null,"source":{"id":"https://openalex.org/S106296714","display_name":"Lecture notes in computer science","issn_l":"0302-9743","issn":["0302-9743","1611-3349"],"is_oa":false,"is_in_doaj":false,"is_indexed_in_scopus":true,"is_core":true,"host_organization":"https://openalex.org/P4310319900","host_organization_name":"Springer Science+Business Media","host_organization_lineage":["https://openalex.org/P4310319965","https://openalex.org/P4310319900"],"host_organization_lineage_names":["Springer Nature","Springer Science+Business Media"],"type":"book series"},"license":null,"license_id":null,"version":null,"is_accepted":false,"is_published":false},{"is_oa":true,"landing_page_url":"http://arxiv.org/abs/1511.06285","pdf_url":"http://arxiv.org/pdf/1511.06285","source":{"id":"https://openalex.org/S4306400194","display_name":"arXiv (Cornell University)","issn_l":null,"issn":null,"is_oa":true,"is_in_doaj":false,"is_indexed_in_scopus":false,"is_core":false,"host_organization":"https://openalex.org/I205783295","host_organization_name":"Cornell University","host_organization_lineage":["https://openalex.org/I205783295"],"host_organization_lineage_names":["Cornell University"],"type":"repository"},"license":null,"license_id":null,"version":"submittedVersion","is_accepted":false,"is_published":false},{"is_oa":true,"landing_page_url":"https://arxiv.org/abs/1511.06285","pdf_url":"https://arxiv.org/pdf/1511.06285","source":{"id":"https://openalex.org/S4306400194","display_name":"arXiv (Cornell University)","issn_l":null,"issn":null,"is_oa":true,"is_in_doaj":false,"is_indexed_in_scopus":false,"is_core":false,"host_organization":"https://openalex.org/I205783295","host_organization_name":"Cornell University","host_organization_lineage":["https://openalex.org/I205783295"],"host_organization_lineage_names":["Cornell University"],"type":"repository"},"license":null,"license_id":null,"version":"submittedVersion","is_accepted":false,"is_published":false},{"is_oa":false,"landing_page_url":"https://api.datacite.org/dois/10.48550/arxiv.1511.06285","pdf_url":null,"source":{"id":"https://openalex.org/S4393179698","display_name":"DataCite API","issn_l":null,"issn":null,"is_oa":true,"is_in_doaj":false,"is_indexed_in_scopus":false,"is_core":false,"host_organization":"https://openalex.org/I4210145204","host_organization_name":"DataCite","host_organization_lineage":["https://openalex.org/I4210145204"],"host_organization_lineage_names":["DataCite"],"type":"metadata"},"license":null,"license_id":null,"version":null}],"best_oa_location":{"is_oa":true,"landing_page_url":"http://arxiv.org/abs/1511.06285","pdf_url":"http://arxiv.org/pdf/1511.06285","source":{"id":"https://openalex.org/S4306400194","display_name":"arXiv (Cornell University)","issn_l":null,"issn":null,"is_oa":true,"is_in_doaj":false,"is_indexed_in_scopus":false,"is_core":false,"host_organization":"https://openalex.org/I205783295","host_organization_name":"Cornell University","host_organization_lineage":["https://openalex.org/I205783295"],"host_organization_lineage_names":["Cornell University"],"type":"repository"},"license":null,"license_id":null,"version":"submittedVersion","is_accepted":false,"is_published":false},"sustainable_development_goals":[{"id":"https://metadata.un.org/sdg/4","score":0.79,"display_name":"Quality education"}],"grants":[],"datasets":[],"versions":["https://openalex.org/W2290619535","https://openalex.org/W3106349295"],"referenced_works_count":17,"referenced_works":["https://openalex.org/W14574270","https://openalex.org/W1543317902","https://openalex.org/W2103610940","https://openalex.org/W2105673178","https://openalex.org/W2107878390","https://openalex.org/W211166471","https://openalex.org/W2139799442","https://openalex.org/W2144600658","https://openalex.org/W2149684865","https://openalex.org/W2184135559","https://openalex.org/W2250852500","https://openalex.org/W2251569308","https://openalex.org/W2405053895","https://openalex.org/W284732977","https://openalex.org/W46893310","https://openalex.org/W630532510","https://openalex.org/W67077090"],"related_works":["https://openalex.org/W4307459710","https://openalex.org/W4293584592","https://openalex.org/W3175595715","https://openalex.org/W3155572818","https://openalex.org/W2986030184","https://openalex.org/W2985215540","https://openalex.org/W2797913374","https://openalex.org/W2786253471","https://openalex.org/W2604275745","https://openalex.org/W2104907655"],"abstract_inverted_index":null,"abstract_inverted_index_v3":null,"cited_by_api_url":"https://api.openalex.org/works?filter=cites:W2290619535","counts_by_year":[{"year":2023,"cited_by_count":1},{"year":2017,"cited_by_count":1}],"updated_date":"2025-03-18T05:25:12.036155","created_date":"2016-06-24"}