Apache Solr Content Extraction Library integrates Apache Tika content extraction framework into Solr
'org.apache.solr:solr-cell:4.10.0'
<dependency>
<groupId>org.apache.solr</groupId>
<artifactId>solr-cell</artifactId>
<version>4.10.0</version>
</dependency>
<dependency org="org.apache.solr" name="solr-cell" rev="4.10.0"/>
"org.apache.solr", "solr-cell", "4.10.0"