2020-11-25 00:24:04 +01:00
|
|
|
use std::convert::TryFrom;
|
2022-02-07 20:48:57 +01:00
|
|
|
use std::fs;
|
2018-10-07 11:54:01 +02:00
|
|
|
use std::fs::File;
|
2020-04-21 21:14:44 +02:00
|
|
|
use std::io::{self, BufRead, BufReader, Read};
|
2021-03-02 09:15:49 +01:00
|
|
|
use std::path::{Path, PathBuf};
|
2018-10-07 11:54:01 +02:00
|
|
|
|
2021-02-27 15:32:07 +01:00
|
|
|
use clircle::{Clircle, Identifier};
|
2018-10-07 16:44:59 +02:00
|
|
|
use content_inspector::{self, ContentType};
|
|
|
|
|
2020-04-22 21:45:47 +02:00
|
|
|
use crate::error::*;
|
2018-10-07 11:54:01 +02:00
|
|
|
|
2020-05-16 01:18:10 +02:00
|
|
|
/// A description of an Input source.
|
2020-05-15 23:21:38 +02:00
|
|
|
/// This tells bat how to refer to the input.
|
2020-05-16 01:18:10 +02:00
|
|
|
#[derive(Clone)]
|
|
|
|
pub struct InputDescription {
|
2020-05-17 00:06:43 +02:00
|
|
|
pub(crate) name: String,
|
|
|
|
|
|
|
|
/// The input title.
|
|
|
|
/// This replaces the name if provided.
|
|
|
|
title: Option<String>,
|
|
|
|
|
|
|
|
/// The input kind.
|
2020-05-16 01:18:10 +02:00
|
|
|
kind: Option<String>,
|
2020-05-17 00:06:43 +02:00
|
|
|
|
|
|
|
/// A summary description of the input.
|
|
|
|
/// Defaults to "{kind} '{name}'"
|
2020-05-16 01:18:10 +02:00
|
|
|
summary: Option<String>,
|
|
|
|
}
|
|
|
|
|
|
|
|
impl InputDescription {
|
|
|
|
/// Creates a description for an input.
|
|
|
|
pub fn new(name: impl Into<String>) -> Self {
|
|
|
|
InputDescription {
|
|
|
|
name: name.into(),
|
2020-05-17 00:06:43 +02:00
|
|
|
title: None,
|
2020-05-16 01:18:10 +02:00
|
|
|
kind: None,
|
|
|
|
summary: None,
|
|
|
|
}
|
|
|
|
}
|
2020-05-15 23:21:38 +02:00
|
|
|
|
2020-05-28 01:26:28 +02:00
|
|
|
pub fn set_kind(&mut self, kind: Option<String>) {
|
2020-05-17 00:06:43 +02:00
|
|
|
self.kind = kind;
|
2020-05-16 01:18:10 +02:00
|
|
|
}
|
|
|
|
|
2020-05-28 01:26:28 +02:00
|
|
|
pub fn set_summary(&mut self, summary: Option<String>) {
|
2020-05-17 00:06:43 +02:00
|
|
|
self.summary = summary;
|
2020-05-16 01:18:10 +02:00
|
|
|
}
|
|
|
|
|
2020-05-28 01:26:28 +02:00
|
|
|
pub fn set_title(&mut self, title: Option<String>) {
|
2020-05-17 00:06:43 +02:00
|
|
|
self.title = title;
|
|
|
|
}
|
|
|
|
|
|
|
|
pub fn title(&self) -> &String {
|
2021-09-10 21:58:46 +02:00
|
|
|
match &self.title {
|
2021-08-18 20:37:38 +02:00
|
|
|
Some(title) => title,
|
2020-05-17 00:06:43 +02:00
|
|
|
None => &self.name,
|
|
|
|
}
|
2020-05-16 01:18:10 +02:00
|
|
|
}
|
2020-05-15 23:21:38 +02:00
|
|
|
|
2020-05-16 01:18:10 +02:00
|
|
|
pub fn kind(&self) -> Option<&String> {
|
|
|
|
self.kind.as_ref()
|
|
|
|
}
|
|
|
|
|
|
|
|
pub fn summary(&self) -> String {
|
|
|
|
self.summary.clone().unwrap_or_else(|| match &self.kind {
|
|
|
|
None => self.name.clone(),
|
|
|
|
Some(kind) => format!("{} '{}'", kind.to_lowercase(), self.name),
|
|
|
|
})
|
|
|
|
}
|
2020-04-21 22:24:47 +02:00
|
|
|
}
|
|
|
|
|
2020-04-22 18:30:06 +02:00
|
|
|
pub(crate) enum InputKind<'a> {
|
2021-03-02 09:15:49 +01:00
|
|
|
OrdinaryFile(PathBuf),
|
2020-04-22 16:27:34 +02:00
|
|
|
StdIn,
|
2020-04-22 18:10:26 +02:00
|
|
|
CustomReader(Box<dyn Read + 'a>),
|
2020-04-22 16:27:34 +02:00
|
|
|
}
|
|
|
|
|
2020-05-16 01:18:10 +02:00
|
|
|
impl<'a> InputKind<'a> {
|
|
|
|
pub fn description(&self) -> InputDescription {
|
|
|
|
match self {
|
2020-05-17 00:06:43 +02:00
|
|
|
InputKind::OrdinaryFile(ref path) => InputDescription::new(path.to_string_lossy()),
|
2020-05-16 01:18:10 +02:00
|
|
|
InputKind::StdIn => InputDescription::new("STDIN"),
|
|
|
|
InputKind::CustomReader(_) => InputDescription::new("READER"),
|
|
|
|
}
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2020-04-22 16:27:34 +02:00
|
|
|
#[derive(Clone, Default)]
|
2020-04-22 18:30:06 +02:00
|
|
|
pub(crate) struct InputMetadata {
|
2021-03-02 09:15:49 +01:00
|
|
|
pub(crate) user_provided_name: Option<PathBuf>,
|
2022-02-07 20:48:57 +01:00
|
|
|
pub(crate) size: Option<u64>,
|
2020-04-22 16:27:34 +02:00
|
|
|
}
|
|
|
|
|
2020-04-22 18:10:26 +02:00
|
|
|
pub struct Input<'a> {
|
2020-04-22 18:30:06 +02:00
|
|
|
pub(crate) kind: InputKind<'a>,
|
|
|
|
pub(crate) metadata: InputMetadata,
|
2020-05-17 00:06:43 +02:00
|
|
|
pub(crate) description: InputDescription,
|
2020-04-22 16:27:34 +02:00
|
|
|
}
|
|
|
|
|
2020-04-22 18:30:06 +02:00
|
|
|
pub(crate) enum OpenedInputKind {
|
2021-03-02 09:15:49 +01:00
|
|
|
OrdinaryFile(PathBuf),
|
2020-04-22 16:27:34 +02:00
|
|
|
StdIn,
|
|
|
|
CustomReader,
|
|
|
|
}
|
|
|
|
|
2020-04-22 18:30:06 +02:00
|
|
|
pub(crate) struct OpenedInput<'a> {
|
|
|
|
pub(crate) kind: OpenedInputKind,
|
|
|
|
pub(crate) metadata: InputMetadata,
|
|
|
|
pub(crate) reader: InputReader<'a>,
|
2020-05-16 01:18:10 +02:00
|
|
|
pub(crate) description: InputDescription,
|
2018-10-07 11:21:41 +02:00
|
|
|
}
|
2018-10-07 11:54:01 +02:00
|
|
|
|
2021-09-26 10:00:40 +02:00
|
|
|
impl OpenedInput<'_> {
|
|
|
|
/// Get the path of the file:
|
|
|
|
/// If this was set by the metadata, that will take priority.
|
|
|
|
/// If it wasn't, it will use the real file path (if available).
|
|
|
|
pub(crate) fn path(&self) -> Option<&PathBuf> {
|
|
|
|
self.metadata
|
|
|
|
.user_provided_name
|
|
|
|
.as_ref()
|
2021-12-08 08:29:10 +01:00
|
|
|
.or(match self.kind {
|
2021-09-26 10:00:40 +02:00
|
|
|
OpenedInputKind::OrdinaryFile(ref path) => Some(path),
|
|
|
|
_ => None,
|
|
|
|
})
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2020-04-22 18:10:26 +02:00
|
|
|
impl<'a> Input<'a> {
|
2021-03-02 09:15:49 +01:00
|
|
|
pub fn ordinary_file(path: impl AsRef<Path>) -> Self {
|
|
|
|
Self::_ordinary_file(path.as_ref())
|
|
|
|
}
|
2021-03-07 14:41:52 +01:00
|
|
|
|
2021-03-02 09:15:49 +01:00
|
|
|
fn _ordinary_file(path: &Path) -> Self {
|
|
|
|
let kind = InputKind::OrdinaryFile(path.to_path_buf());
|
2022-02-07 20:48:57 +01:00
|
|
|
let metadata = InputMetadata {
|
|
|
|
size: fs::metadata(path).map(|m| m.len()).ok(),
|
|
|
|
..InputMetadata::default()
|
|
|
|
};
|
|
|
|
|
2020-04-22 16:27:34 +02:00
|
|
|
Input {
|
2020-05-17 00:06:43 +02:00
|
|
|
description: kind.description(),
|
2022-02-07 20:48:57 +01:00
|
|
|
metadata,
|
2020-05-17 00:06:43 +02:00
|
|
|
kind,
|
2018-10-07 11:54:01 +02:00
|
|
|
}
|
|
|
|
}
|
2020-04-21 22:24:47 +02:00
|
|
|
|
2020-04-22 16:27:34 +02:00
|
|
|
pub fn stdin() -> Self {
|
2020-05-17 00:06:43 +02:00
|
|
|
let kind = InputKind::StdIn;
|
2020-04-22 16:27:34 +02:00
|
|
|
Input {
|
2020-05-17 00:06:43 +02:00
|
|
|
description: kind.description(),
|
2020-04-22 16:27:34 +02:00
|
|
|
metadata: InputMetadata::default(),
|
2020-05-17 00:06:43 +02:00
|
|
|
kind,
|
2020-04-22 16:27:34 +02:00
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2020-04-22 18:10:26 +02:00
|
|
|
pub fn from_reader(reader: Box<dyn Read + 'a>) -> Self {
|
2020-05-17 00:06:43 +02:00
|
|
|
let kind = InputKind::CustomReader(reader);
|
2020-04-22 18:10:26 +02:00
|
|
|
Input {
|
2020-05-17 00:06:43 +02:00
|
|
|
description: kind.description(),
|
2020-04-22 18:10:26 +02:00
|
|
|
metadata: InputMetadata::default(),
|
2020-05-17 00:06:43 +02:00
|
|
|
kind,
|
2020-04-22 18:10:26 +02:00
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2020-04-22 16:27:34 +02:00
|
|
|
pub fn is_stdin(&self) -> bool {
|
2021-01-09 15:27:54 +01:00
|
|
|
matches!(self.kind, InputKind::StdIn)
|
2020-04-22 16:27:34 +02:00
|
|
|
}
|
|
|
|
|
2021-03-07 14:41:52 +01:00
|
|
|
pub fn with_name(self, provided_name: Option<impl AsRef<Path>>) -> Self {
|
2021-03-02 09:15:49 +01:00
|
|
|
self._with_name(provided_name.as_ref().map(|it| it.as_ref()))
|
|
|
|
}
|
2021-03-07 14:41:52 +01:00
|
|
|
|
|
|
|
fn _with_name(mut self, provided_name: Option<&Path>) -> Self {
|
2020-11-05 22:29:04 +01:00
|
|
|
if let Some(name) = provided_name {
|
|
|
|
self.description.name = name.to_string_lossy().to_string()
|
2020-05-17 00:06:43 +02:00
|
|
|
}
|
|
|
|
|
2020-04-22 16:27:34 +02:00
|
|
|
self.metadata.user_provided_name = provided_name.map(|n| n.to_owned());
|
2020-04-22 22:41:25 +02:00
|
|
|
self
|
2020-04-22 16:27:34 +02:00
|
|
|
}
|
|
|
|
|
2020-05-17 00:06:43 +02:00
|
|
|
pub fn description(&self) -> &InputDescription {
|
|
|
|
&self.description
|
|
|
|
}
|
|
|
|
|
|
|
|
pub fn description_mut(&mut self) -> &mut InputDescription {
|
|
|
|
&mut self.description
|
2020-05-16 01:18:10 +02:00
|
|
|
}
|
|
|
|
|
2021-02-27 15:32:07 +01:00
|
|
|
pub(crate) fn open<R: BufRead + 'a>(
|
|
|
|
self,
|
|
|
|
stdin: R,
|
|
|
|
stdout_identifier: Option<&Identifier>,
|
|
|
|
) -> Result<OpenedInput<'a>> {
|
2020-05-16 01:18:10 +02:00
|
|
|
let description = self.description().clone();
|
2020-04-22 16:27:34 +02:00
|
|
|
match self.kind {
|
2021-02-27 15:32:07 +01:00
|
|
|
InputKind::StdIn => {
|
|
|
|
if let Some(stdout) = stdout_identifier {
|
|
|
|
let input_identifier = Identifier::try_from(clircle::Stdio::Stdin)
|
|
|
|
.map_err(|e| format!("Stdin: Error identifying file: {}", e))?;
|
|
|
|
if stdout.surely_conflicts_with(&input_identifier) {
|
|
|
|
return Err("IO circle detected. The input from stdin is also an output. Aborting to avoid infinite loop.".into());
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
|
|
|
Ok(OpenedInput {
|
|
|
|
kind: OpenedInputKind::StdIn,
|
|
|
|
description,
|
|
|
|
metadata: self.metadata,
|
|
|
|
reader: InputReader::new(stdin),
|
|
|
|
})
|
|
|
|
}
|
|
|
|
|
2020-04-22 16:27:34 +02:00
|
|
|
InputKind::OrdinaryFile(path) => Ok(OpenedInput {
|
|
|
|
kind: OpenedInputKind::OrdinaryFile(path.clone()),
|
2020-05-16 01:18:10 +02:00
|
|
|
description,
|
2020-04-22 16:27:34 +02:00
|
|
|
metadata: self.metadata,
|
|
|
|
reader: {
|
2021-02-27 15:32:07 +01:00
|
|
|
let mut file = File::open(&path)
|
2020-04-22 16:27:34 +02:00
|
|
|
.map_err(|e| format!("'{}': {}", path.to_string_lossy(), e))?;
|
|
|
|
if file.metadata()?.is_dir() {
|
|
|
|
return Err(format!("'{}' is a directory.", path.to_string_lossy()).into());
|
|
|
|
}
|
2021-02-27 15:32:07 +01:00
|
|
|
|
|
|
|
if let Some(stdout) = stdout_identifier {
|
|
|
|
let input_identifier = Identifier::try_from(file).map_err(|e| {
|
|
|
|
format!("{}: Error identifying file: {}", path.to_string_lossy(), e)
|
|
|
|
})?;
|
|
|
|
if stdout.surely_conflicts_with(&input_identifier) {
|
|
|
|
return Err(format!(
|
|
|
|
"IO circle detected. The input from '{}' is also an output. Aborting to avoid infinite loop.",
|
|
|
|
path.to_string_lossy()
|
|
|
|
)
|
|
|
|
.into());
|
|
|
|
}
|
2021-06-13 21:32:08 +02:00
|
|
|
file = input_identifier.into_inner().expect("The file was lost in the clircle::Identifier, this should not have happened...");
|
2021-02-27 15:32:07 +01:00
|
|
|
}
|
|
|
|
|
2020-04-22 16:27:34 +02:00
|
|
|
InputReader::new(BufReader::new(file))
|
|
|
|
},
|
|
|
|
}),
|
|
|
|
InputKind::CustomReader(reader) => Ok(OpenedInput {
|
2020-05-16 01:18:10 +02:00
|
|
|
description,
|
2020-04-22 16:27:34 +02:00
|
|
|
kind: OpenedInputKind::CustomReader,
|
|
|
|
metadata: self.metadata,
|
|
|
|
reader: InputReader::new(BufReader::new(reader)),
|
|
|
|
}),
|
|
|
|
}
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2020-04-22 18:30:06 +02:00
|
|
|
pub(crate) struct InputReader<'a> {
|
2020-04-21 22:24:47 +02:00
|
|
|
inner: Box<dyn BufRead + 'a>,
|
|
|
|
pub(crate) first_line: Vec<u8>,
|
|
|
|
pub(crate) content_type: Option<ContentType>,
|
|
|
|
}
|
|
|
|
|
|
|
|
impl<'a> InputReader<'a> {
|
|
|
|
fn new<R: BufRead + 'a>(mut reader: R) -> InputReader<'a> {
|
|
|
|
let mut first_line = vec![];
|
|
|
|
reader.read_until(b'\n', &mut first_line).ok();
|
|
|
|
|
|
|
|
let content_type = if first_line.is_empty() {
|
|
|
|
None
|
|
|
|
} else {
|
|
|
|
Some(content_inspector::inspect(&first_line[..]))
|
|
|
|
};
|
|
|
|
|
|
|
|
if content_type == Some(ContentType::UTF_16LE) {
|
|
|
|
reader.read_until(0x00, &mut first_line).ok();
|
|
|
|
}
|
|
|
|
|
|
|
|
InputReader {
|
|
|
|
inner: Box::new(reader),
|
|
|
|
first_line,
|
|
|
|
content_type,
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
|
|
|
pub(crate) fn read_line(&mut self, buf: &mut Vec<u8>) -> io::Result<bool> {
|
2021-09-10 21:52:09 +02:00
|
|
|
if !self.first_line.is_empty() {
|
|
|
|
buf.append(&mut self.first_line);
|
|
|
|
return Ok(true);
|
|
|
|
}
|
2020-04-21 22:24:47 +02:00
|
|
|
|
2021-09-10 21:52:09 +02:00
|
|
|
let res = self.inner.read_until(b'\n', buf).map(|size| size > 0)?;
|
2020-04-21 22:24:47 +02:00
|
|
|
|
2021-09-10 21:52:09 +02:00
|
|
|
if self.content_type == Some(ContentType::UTF_16LE) {
|
|
|
|
let _ = self.inner.read_until(0x00, buf);
|
2020-04-21 22:24:47 +02:00
|
|
|
}
|
2021-09-10 21:52:09 +02:00
|
|
|
|
|
|
|
Ok(res)
|
2020-04-21 22:24:47 +02:00
|
|
|
}
|
2018-10-07 11:54:01 +02:00
|
|
|
}
|
2018-10-07 12:29:38 +02:00
|
|
|
|
|
|
|
#[test]
|
|
|
|
fn basic() {
|
2018-10-07 13:25:49 +02:00
|
|
|
let content = b"#!/bin/bash\necho hello";
|
2020-04-21 21:19:06 +02:00
|
|
|
let mut reader = InputReader::new(&content[..]);
|
2018-10-07 12:29:38 +02:00
|
|
|
|
2018-10-07 13:47:54 +02:00
|
|
|
assert_eq!(b"#!/bin/bash\n", &reader.first_line[..]);
|
2018-10-07 13:25:49 +02:00
|
|
|
|
2018-10-07 12:29:38 +02:00
|
|
|
let mut buffer = vec![];
|
|
|
|
|
|
|
|
let res = reader.read_line(&mut buffer);
|
|
|
|
assert!(res.is_ok());
|
2021-08-02 15:06:57 +02:00
|
|
|
assert!(res.unwrap());
|
2018-10-07 13:25:49 +02:00
|
|
|
assert_eq!(b"#!/bin/bash\n", &buffer[..]);
|
2018-10-07 12:29:38 +02:00
|
|
|
|
|
|
|
buffer.clear();
|
|
|
|
|
|
|
|
let res = reader.read_line(&mut buffer);
|
|
|
|
assert!(res.is_ok());
|
2021-08-02 15:06:57 +02:00
|
|
|
assert!(res.unwrap());
|
2018-10-07 13:25:49 +02:00
|
|
|
assert_eq!(b"echo hello", &buffer[..]);
|
2018-10-07 12:29:38 +02:00
|
|
|
|
|
|
|
buffer.clear();
|
|
|
|
|
|
|
|
let res = reader.read_line(&mut buffer);
|
|
|
|
assert!(res.is_ok());
|
2021-08-02 15:06:57 +02:00
|
|
|
assert!(!res.unwrap());
|
2018-10-07 12:29:38 +02:00
|
|
|
assert!(buffer.is_empty());
|
|
|
|
}
|
2018-10-07 16:44:59 +02:00
|
|
|
|
|
|
|
#[test]
|
|
|
|
fn utf16le() {
|
|
|
|
let content = b"\xFF\xFE\x73\x00\x0A\x00\x64\x00";
|
2020-04-21 21:19:06 +02:00
|
|
|
let mut reader = InputReader::new(&content[..]);
|
2018-10-07 16:44:59 +02:00
|
|
|
|
|
|
|
assert_eq!(b"\xFF\xFE\x73\x00\x0A\x00", &reader.first_line[..]);
|
|
|
|
|
|
|
|
let mut buffer = vec![];
|
|
|
|
|
|
|
|
let res = reader.read_line(&mut buffer);
|
|
|
|
assert!(res.is_ok());
|
2021-08-02 15:06:57 +02:00
|
|
|
assert!(res.unwrap());
|
2018-10-07 16:44:59 +02:00
|
|
|
assert_eq!(b"\xFF\xFE\x73\x00\x0A\x00", &buffer[..]);
|
|
|
|
|
|
|
|
buffer.clear();
|
|
|
|
|
|
|
|
let res = reader.read_line(&mut buffer);
|
|
|
|
assert!(res.is_ok());
|
2021-08-02 15:06:57 +02:00
|
|
|
assert!(res.unwrap());
|
2018-10-07 16:44:59 +02:00
|
|
|
assert_eq!(b"\x64\x00", &buffer[..]);
|
|
|
|
|
|
|
|
buffer.clear();
|
|
|
|
|
|
|
|
let res = reader.read_line(&mut buffer);
|
|
|
|
assert!(res.is_ok());
|
2021-08-02 15:06:57 +02:00
|
|
|
assert!(!res.unwrap());
|
2018-10-07 16:44:59 +02:00
|
|
|
assert!(buffer.is_empty());
|
|
|
|
}
|