Always show summary of reformatting

[etc/vim.git] / black.py
diff --git a/black.py b/black.py

index 730c64de0b1883ff42e906d90d6985d4ddec9b51..19a023cff52db2e92d5cc0950125e20c96d24dff 100644 (file)
--- a/black.py
+++ b/black.py
@@ -4,6 +4,7 @@ from asyncio.base_events import BaseEventLoop
  from concurrent.futures import Executor, ProcessPoolExecutor
  from enum import Enum, Flag
  from functools import partial, wraps
+import io
  import keyword
  import logging
  from multiprocessing import Manager
@@ -119,6 +120,13 @@ class WriteBack(Enum):
      YES = 1
      DIFF = 2
  
+    @classmethod
+    def from_configuration(cls, *, check: bool, diff: bool) -> "WriteBack":
+        if check and not diff:
+            return cls.NO
+
+        return cls.DIFF if diff else cls.YES
+
  
  class Changed(Enum):
      NO = 0
@@ -132,6 +140,19 @@ class FileMode(Flag):
      PYI = 2
      NO_STRING_NORMALIZATION = 4
  
+    @classmethod
+    def from_configuration(
+        cls, *, py36: bool, pyi: bool, skip_string_normalization: bool
+    ) -> "FileMode":
+        mode = cls.AUTO_DETECT
+        if py36:
+            mode |= cls.PYTHON36
+        if pyi:
+            mode |= cls.PYI
+        if skip_string_normalization:
+            mode |= cls.NO_STRING_NORMALIZATION
+        return mode
+
  
  @click.command()
  @click.option(
@@ -218,6 +239,15 @@ class FileMode(Flag):
          "silence those with 2>/dev/null."
      ),
  )
+@click.option(
+    "-v",
+    "--verbose",
+    is_flag=True,
+    help=(
+        "Also emit messages to stderr about files that were not changed or were "
+        "ignored due to --exclude=."
+    ),
+)
  @click.version_option(version=__version__)
  @click.argument(
      "src",
@@ -237,12 +267,18 @@ def main(
      py36: bool,
      skip_string_normalization: bool,
      quiet: bool,
+    verbose: bool,
      include: str,
      exclude: str,
      src: List[str],
  ) -> None:
      """The uncompromising code formatter."""
-    sources: List[Path] = []
+    write_back = WriteBack.from_configuration(check=check, diff=diff)
+    mode = FileMode.from_configuration(
+        py36=py36, pyi=pyi, skip_string_normalization=skip_string_normalization
+    )
+    report = Report(check=check, quiet=quiet, verbose=verbose)
+    sources: Set[Path] = set()
      try:
          include_regex = re.compile(include)
      except re.error:
@@ -257,39 +293,23 @@ def main(
      for s in src:
          p = Path(s)
          if p.is_dir():
-            sources.extend(
-                gen_python_files_in_dir(p, root, include_regex, exclude_regex)
+            sources.update(
+                gen_python_files_in_dir(p, root, include_regex, exclude_regex, report)
              )
-        elif p.is_file():
+        elif p.is_file() or s == "-":
              # if a file was explicitly given, we don't care about its extension
-            sources.append(p)
-        elif s == "-":
-            sources.append(Path("-"))
+            sources.add(p)
          else:
              err(f"invalid path: {s}")
-
-    if check and not diff:
-        write_back = WriteBack.NO
-    elif diff:
-        write_back = WriteBack.DIFF
-    else:
-        write_back = WriteBack.YES
-    mode = FileMode.AUTO_DETECT
-    if py36:
-        mode |= FileMode.PYTHON36
-    if pyi:
-        mode |= FileMode.PYI
-    if skip_string_normalization:
-        mode |= FileMode.NO_STRING_NORMALIZATION
-    report = Report(check=check, quiet=quiet)
      if len(sources) == 0:
-        out("No paths given. Nothing to do 😴")
+        if verbose or not quiet:
+            out("No paths given. Nothing to do 😴")
          ctx.exit(0)
          return
  
      elif len(sources) == 1:
          reformat_one(
-            src=sources[0],
+            src=sources.pop(),
              line_length=line_length,
              fast=fast,
              write_back=write_back,
@@ -314,9 +334,9 @@ def main(
              )
          finally:
              shutdown(loop)
-        if not quiet:
-            out("All done! ✨ 🍰 ✨")
-            click.echo(str(report))
+    if verbose or not quiet:
+        out("All done! ✨ 🍰 ✨")
+        click.echo(str(report))
      ctx.exit(report.return_code)
  
  
@@ -364,7 +384,7 @@ def reformat_one(
  
  
  async def schedule_formatting(
-    sources: List[Path],
+    sources: Set[Path],
      line_length: int,
      fast: bool,
      write_back: WriteBack,
@@ -384,7 +404,7 @@ async def schedule_formatting(
      if write_back != WriteBack.DIFF:
          cache = read_cache(line_length, mode)
          sources, cached = filter_cached(cache, sources)
-        for src in cached:
+        for src in sorted(cached):
              report.done(src, Changed.CACHED)
      cancelled = []
      formatted = []
@@ -447,8 +467,9 @@ def format_file_in_place(
      """
      if src.suffix == ".pyi":
          mode |= FileMode.PYI
-    with tokenize.open(src) as src_buffer:
-        src_contents = src_buffer.read()
+
+    with open(src, "rb") as buf:
+        newline, encoding, src_contents = prepare_input(buf.read())
      try:
          dst_contents = format_file_contents(
              src_contents, line_length=line_length, fast=fast, mode=mode
@@ -457,7 +478,7 @@ def format_file_in_place(
          return False
  
      if write_back == write_back.YES:
-        with open(src, "w", encoding=src_buffer.encoding) as f:
+        with open(src, "w", encoding=encoding, newline=newline) as f:
              f.write(dst_contents)
      elif write_back == write_back.DIFF:
          src_name = f"{src}  (original)"
@@ -466,7 +487,14 @@ def format_file_in_place(
          if lock:
              lock.acquire()
          try:
-            sys.stdout.write(diff_contents)
+            f = io.TextIOWrapper(
+                sys.stdout.buffer,
+                encoding=encoding,
+                newline=newline,
+                write_through=True,
+            )
+            f.write(diff_contents)
+            f.detach()
          finally:
              if lock:
                  lock.release()
@@ -485,7 +513,7 @@ def format_stdin_to_stdout(
      `line_length`, `fast`, `is_pyi`, and `force_py36` arguments are passed to
      :func:`format_file_contents`.
      """
-    src = sys.stdin.read()
+    newline, encoding, src = prepare_input(sys.stdin.buffer.read())
      dst = src
      try:
          dst = format_file_contents(src, line_length=line_length, fast=fast, mode=mode)
@@ -496,11 +524,25 @@ def format_stdin_to_stdout(
  
      finally:
          if write_back == WriteBack.YES:
-            sys.stdout.write(dst)
+            f = io.TextIOWrapper(
+                sys.stdout.buffer,
+                encoding=encoding,
+                newline=newline,
+                write_through=True,
+            )
+            f.write(dst)
+            f.detach()
          elif write_back == WriteBack.DIFF:
              src_name = "<stdin>  (original)"
              dst_name = "<stdin>  (formatted)"
-            sys.stdout.write(diff(src, dst, src_name, dst_name))
+            f = io.TextIOWrapper(
+                sys.stdout.buffer,
+                encoding=encoding,
+                newline=newline,
+                write_through=True,
+            )
+            f.write(diff(src, dst, src_name, dst_name))
+            f.detach()
  
  
  def format_file_contents(
@@ -561,6 +603,19 @@ def format_str(
      return dst_contents
  
  
+def prepare_input(src: bytes) -> Tuple[str, str, str]:
+    """Analyze `src` and return a tuple of (newline, encoding, decoded_contents)
+
+    Where `newline` is either CRLF or LF, and `decoded_contents` is decoded with
+    universal newlines (i.e. only LF).
+    """
+    srcbuf = io.BytesIO(src)
+    encoding, lines = tokenize.detect_encoding(srcbuf.readline)
+    newline = "\r\n" if b"\r\n" == lines[0][-2:] else "\n"
+    srcbuf.seek(0)
+    return newline, encoding, io.TextIOWrapper(srcbuf, encoding).read()
+
+
  GRAMMARS = [
      pygram.python_grammar_no_print_statement_no_exec_statement,
      pygram.python_grammar_no_print_statement,
@@ -572,8 +627,7 @@ def lib2to3_parse(src_txt: str) -> Node:
      """Given a string with source, return the lib2to3 Node."""
      grammar = pygram.python_grammar_no_print_statement
      if src_txt[-1] != "\n":
-        nl = "\r\n" if "\r\n" in src_txt[:1024] else "\n"
-        src_txt += nl
+        src_txt += "\n"
      for grammar in GRAMMARS:
          drv = driver.Driver(grammar, pytree.convert)
          try:
@@ -2795,22 +2849,29 @@ def get_future_imports(node: Node) -> Set[str]:
  
  
  def gen_python_files_in_dir(
-    path: Path, root: Path, include: Pattern[str], exclude: Pattern[str]
+    path: Path,
+    root: Path,
+    include: Pattern[str],
+    exclude: Pattern[str],
+    report: "Report",
  ) -> Iterator[Path]:
      """Generate all files under `path` whose paths are not excluded by the
      `exclude` regex, but are included by the `include` regex.
+
+    `report` is where output about exclusions goes.
      """
      assert root.is_absolute(), f"INTERNAL ERROR: `root` must be absolute but is {root}"
      for child in path.iterdir():
-        normalized_path = child.resolve().relative_to(root).as_posix()
+        normalized_path = "/" + child.resolve().relative_to(root).as_posix()
          if child.is_dir():
              normalized_path += "/"
          exclude_match = exclude.search(normalized_path)
          if exclude_match and exclude_match.group(0):
+            report.path_ignored(child, f"matches --exclude={exclude.pattern}")
              continue
  
          if child.is_dir():
-            yield from gen_python_files_in_dir(child, root, include, exclude)
+            yield from gen_python_files_in_dir(child, root, include, exclude, report)
  
          elif child.is_file():
              include_match = include.search(normalized_path)
@@ -2853,6 +2914,7 @@ class Report:
  
      check: bool = False
      quiet: bool = False
+    verbose: bool = False
      change_count: int = 0
      same_count: int = 0
      failure_count: int = 0
@@ -2861,11 +2923,11 @@ class Report:
          """Increment the counter for successful reformatting. Write out a message."""
          if changed is Changed.YES:
              reformatted = "would reformat" if self.check else "reformatted"
-            if not self.quiet:
+            if self.verbose or not self.quiet:
                  out(f"{reformatted} {src}")
              self.change_count += 1
          else:
-            if not self.quiet:
+            if self.verbose:
                  if changed is Changed.NO:
                      msg = f"{src} already well formatted, good job."
                  else:
@@ -2878,6 +2940,10 @@ class Report:
          err(f"error: cannot format {src}: {message}")
          self.failure_count += 1
  
+    def path_ignored(self, path: Path, message: str) -> None:
+        if self.verbose:
+            out(f"{path} ignored: {message}", bold=False)
+
      @property
      def return_code(self) -> int:
          """Return the exit code that the app should use.
@@ -3238,26 +3304,24 @@ def get_cache_info(path: Path) -> CacheInfo:
      return stat.st_mtime, stat.st_size
  
  
-def filter_cached(
-    cache: Cache, sources: Iterable[Path]
-) -> Tuple[List[Path], List[Path]]:
-    """Split a list of paths into two.
+def filter_cached(cache: Cache, sources: Iterable[Path]) -> Tuple[Set[Path], Set[Path]]:
+    """Split an iterable of paths in `sources` into two sets.
  
-    The first list contains paths of files that modified on disk or are not in the
-    cache. The other list contains paths to non-modified files.
+    The first contains paths of files that modified on disk or are not in the
+    cache. The other contains paths to non-modified files.
      """
-    todo, done = [], []
+    todo, done = set(), set()
      for src in sources:
          src = src.resolve()
          if cache.get(src) != get_cache_info(src):
-            todo.append(src)
+            todo.add(src)
          else:
-            done.append(src)
+            done.add(src)
      return todo, done
  
  
  def write_cache(
-    cache: Cache, sources: List[Path], line_length: int, mode: FileMode
+    cache: Cache, sources: Iterable[Path], line_length: int, mode: FileMode
  ) -> None:
      """Update the cache file."""
      cache_file = get_cache_file(line_length, mode)