java将word转换为html

	public static void main(String[] args) throws Exception {
		String filePath = "C:/Users/Administrator/Desktop/92个诊疗方案及临床路径/";
		File file = new File(filePath);
		File[] files = file.listFiles();
		String name = null;
		for (File file2 : files) {
			Thread.sleep(500);
			name = file2.getName().substring(0, file2.getName().lastIndexOf("."));
			System.out.println(file2.getName());
			if (file2.getName().endsWith(".docx") || file2.getName().endsWith(".DOCX")) {
				CaseHtm.docx(filePath ,file2.getName(),name +".htm");
			}else{
				CaseHtm.dox(filePath ,file2.getName(),name +".htm");
			}
			
        }
	}
	/**
	 * 转换docx
	 * @param filePath
	 * @param fileName
	 * @param htmlName
	 * @throws Exception
	 */
	public static void docx(String filePath ,String fileName,String htmlName) throws Exception{
		final String file = filePath + fileName;
		File f = new File(file);  
		// ) 加载word文档生成 XWPFDocument对象
		InputStream in = new FileInputStream(f);
		XWPFDocument document = new XWPFDocument(in);
		// ) 解析 XHTML配置 (这里设置IURIResolver来设置图片存放的目录)
		File imageFolderFile = new File(filePath);
		XHTMLOptions options = XHTMLOptions.create().URIResolver(new FileURIResolver(imageFolderFile));
		options.setExtractor(new FileImageExtractor(imageFolderFile));
		options.setIgnoreStylesIfUnused(false);
		options.setFragment(true);
		// ) 将 XWPFDocument转换成XHTML
		OutputStream out = new FileOutputStream(new File(filePath + htmlName));
		XHTMLConverter.getInstance().convert(document, out, options);
	}
	/**
	 * 转换doc
	 * @param filePath
	 * @param fileName
	 * @param htmlName
	 * @throws Exception
	 */
	public static void dox(String filePath ,String fileName,String htmlName) throws Exception{
        final String file = filePath + fileName;
        InputStream input = new FileInputStream(new File(file));
        HWPFDocument wordDocument = new HWPFDocument(input);
        WordToHtmlConverter wordToHtmlConverter = new WordToHtmlConverter(DocumentBuilderFactory.newInstance().newDocumentBuilder().newDocument());
       //解析word文档
       wordToHtmlConverter.processDocument(wordDocument);
       Document htmlDocument = wordToHtmlConverter.getDocument();
       
       File htmlFile = new File(filePath + htmlName);
       OutputStream outStream = new FileOutputStream(htmlFile);

       DOMSource domSource = new DOMSource(htmlDocument);
       StreamResult streamResult = new StreamResult(outStream);

       TransformerFactory factory = TransformerFactory.newInstance();
       Transformer serializer = factory.newTransformer();
       serializer.setOutputProperty(OutputKeys.ENCODING, "utf-8");
       serializer.setOutputProperty(OutputKeys.INDENT, "yes");
       serializer.setOutputProperty(OutputKeys.METHOD, "html");
       
       serializer.transform(domSource, streamResult);
       outStream.close();
   }


<dependency>
    <groupId>fr.opensagres.xdocreport</groupId>
    <artifactId>fr.opensagres.xdocreport.document</artifactId>
    <version>1.0.5</version>
</dependency>
<dependency>  
    <groupId>fr.opensagres.xdocreport</groupId>  
    <artifactId>org.apache.poi.xwpf.converter.xhtml</artifactId>  
    <version>1.0.5</version>  
</dependency>
    <dependency>
    <groupId>org.apache.poi</groupId>
    <artifactId>poi</artifactId>
    <version>3.12</version>
</dependency>
<dependency>
    <groupId>org.apache.poi</groupId>
    <artifactId>poi-scratchpad</artifactId>
    <version>3.12</version>
</dependency>



  • 1
    点赞
  • 15
    收藏
    觉得还不错? 一键收藏
  • 12
    评论
Java中将Word转换HTML需要使用Apache POI和JodConverter库。而将HTML导入到UEditor编辑器中,需要在前端使用UEditor富文本编辑器,并在后端使用Java Web框架(如Spring MVC)来处理上传文件和保存HTML内容。 以下是一个简单的Java代码示例,演示如何将Word文档转换HTML,并将HTML保存到服务器上的文件中: ```java import java.io.File; import java.io.FileOutputStream; import java.io.IOException; import org.apache.poi.hwpf.HWPFDocument; import org.jodconverter.JodConverter; import org.jodconverter.office.LocalOfficeManager; import org.jodconverter.office.OfficeException; import org.jodconverter.office.OfficeManager; public class WordToHtmlConverter { public static void main(String[] args) { // Start the office manager OfficeManager officeManager = LocalOfficeManager.builder().officeHome("/usr/lib/libreoffice").install().build(); try { officeManager.start(); // Load the Word document HWPFDocument document = new HWPFDocument(new FileInputStream("document.doc")); // Create a temporary file for the HTML output File tempFile = File.createTempFile("document", ".html"); // Convert the document to HTML JodConverter.convert(document).to(tempFile).execute(); // Save the HTML content to a file on the server FileOutputStream fos = new FileOutputStream("output.html"); fos.write(Files.readAllBytes(tempFile.toPath())); fos.close(); // Delete the temporary file tempFile.delete(); } catch (IOException | OfficeException e) { e.printStackTrace(); } finally { // Stop the office manager officeManager.stop(); } } } ``` 在上述代码中,我们使用JodConverter将Word文档转换HTML,并将HTML内容保存到服务器上的文件中。然后,我们可以在Java Web应用程序中使用Spring MVC等框架来处理上传文件,并将HTML内容保存到数据库中。最后,在前端使用UEditor富文本编辑器时,我们可以将HTML内容加载到编辑器中,让用户进行编辑和保存。
评论 12
添加红包

请填写红包祝福语或标题

红包个数最小为10个

红包金额最低5元

当前余额3.43前往充值 >
需支付:10.00
成就一亿技术人!
领取后你会自动成为博主和红包主的粉丝 规则
hope_wisdom
发出的红包
实付
使用余额支付
点击重新获取
扫码支付
钱包余额 0

抵扣说明:

1.余额是钱包充值的虚拟货币,按照1:1的比例进行支付金额的抵扣。
2.余额无法直接购买下载,可以购买VIP、付费专栏及课程。

余额充值