@inproceedings{6e1028b3d2334075addcbb49b0c26875,
title = "The hidden Web, XML and the Semantic Web: Scientific data management perspectives",
abstract = "The World Wide Web no longer consists just of HTML pages. Our work sheds light on a number of trends on the Internet that go beyond simple Web pages. The hidden Web provides a wealth of data in semi-structured form, accessible through Web forms and Web services. These services, as well as numerous other applications on the Web, commonly use XML, the eXtensible Markup Language. XML has become the lingua franca of the Internet that allows customized markups to be defined for specific domains. On top of XML, the Semantic Web grows as a common structured data source. In this work, we first explain each of these developments in detail. Using real-world examples from scientific domains of great interest today, we then demonstrate how these new developments can assist the managing, harvesting, and organization of data on the Web. On the way, we also illustrate the current research avenues in these domains. We believe that this effort would help bridge multiple database tracks, thereby attracting researchers with a view to extend database technology.",
keywords = "Deep web, Domain-specific markup languages, Hidden Web, Multidisciplinary work, Scientific data, Semantic web, XML",
author = "Suchanek, \{Fabian M.\} and Varde, \{Aparna S.\} and Richi Nayak and Pierre Senellart",
year = "2011",
month = jan,
day = "1",
doi = "10.1145/1951365.1951433",
language = "English",
isbn = "9781450305280",
series = "ACM International Conference Proceeding Series",
publisher = "Association for Computing Machinery",
pages = "534--537",
booktitle = "Advances in Database Technology - EDBT 2011",
note = "14th International Conference on Extending Database Technology: Advances in Database Technology, EDBT 2011 ; Conference date: 22-03-2011 Through 24-03-2011",
}