Apache Tika is a toolkit for detecting and extracting metadata and structured text content from various documents using existing parser libraries.
'org.apache.tika:tika-parent:0.8'
<dependency>
<groupId>org.apache.tika</groupId>
<artifactId>tika-parent</artifactId>
<version>0.8</version>
</dependency>
<dependency org="org.apache.tika" name="tika-parent" rev="0.8"/>
"org.apache.tika", "tika-parent", "0.8"